@misc{indiciae981297781585, title = {Efficient Arbitrary Precision Acceleration for Large Language Models on GPU Tensor Cores}, author = {Shaobo Ma and Chao Fang and Haikuo Shao and Zhongfeng Wang}, year = {2024}, doi = {10.1145/3658617.3697668}, url = {https://arxiv.org/abs/2409.17870}, note = {Source identifier: 2409.17870} }