(self, weight)
| 82 | stochastic_rounding=0, |
| 83 | ) |
| 84 | self.extra["convrot_groupsize"] = torch.tensor(convrot_groupsize, dtype=torch.int32, device=w.device) |
| 85 | return w_q.contiguous(), params.scale.to(torch.float32).contiguous(), self.extra |
| 86 | |
| 87 | |
| 88 | @CONVERT_WEIGHT_REGISTER("fp8") |
| 89 | class QuantWeightFP8(QuantTemplate): |