@misc{indiciae28f949dd7384, title = {Scalable Efficient Training of Large Language Models with Low-dimensional Projected Attention}, author = {Xingtai Lv and Ning Ding and Kaiyan Zhang and Ermo Hua and Ganqu Cui and Bowen Zhou}, year = {2024}, url = {https://arxiv.org/abs/2411.02063}, note = {Source identifier: 2411.02063} }