Merge commit '8c1a95991330118930f23e6a2ba8e76068d8ca22' into develop

This commit is contained in:
assistant-librarian[bot]
2025-09-25 18:15:45 +00:00
parent 0e513e86a4
commit b8448ab68d
51 changed files with 281 additions and 236 deletions

View File

@@ -598,8 +598,8 @@ struct FlatmmKernel
CK_TILE_DEVICE void operator()(KernelArgs kargs) const
{
const auto [iM, iN] = TilePartitioner{kargs.M, kargs.N}.GetOutputTileIndex(blockIdx.x);
const index_t i_m = __builtin_amdgcn_readfirstlane(iM * TilePartitioner::MPerBlock);
const index_t i_n = __builtin_amdgcn_readfirstlane(iN * TilePartitioner::NPerBlock);
const index_t i_m = amd_wave_read_first_lane(iM * TilePartitioner::MPerBlock);
const index_t i_n = amd_wave_read_first_lane(iN * TilePartitioner::NPerBlock);
const SplitKBatchOffset splitk_batch_offset(kargs);
// options