@misc{indiciae32a27ac74512, title = {VQ-Logits: Compressing the Output Bottleneck of Large Language Models via Vector Quantized Logits}, author = {Jintian Shao and Hongyi Huang and Jiayi Wu and YiMing Cheng and ZhiYu Wu and You Shan and MingKai Zheng}, year = {2026}, url = {https://arxiv.org/abs/2505.10202}, note = {Source identifier: 2505.10202} }