(A: torch.Tensor, B: torch.Tensor)
| 77 | |
| 78 | @register_kernel("bitsandbytes::int8_linear_matmul", "cuda") |
| 79 | def _(A: torch.Tensor, B: torch.Tensor): |
| 80 | out = torch.empty((*A.shape[:-1], B.shape[0]), device=A.device, dtype=torch.int32) |
| 81 | return _int8_linear_matmul_impl(A, B, out) |
| 82 | |
| 83 | |
| 84 | @register_kernel("bitsandbytes::int8_linear_matmul.out", "cuda") |
nothing calls this directly
no test coverage detected