@misc{indiciaea98cc4f8e565, title = {When Linear Attention Meets Autoregressive Decoding: Towards More Effective and Efficient Linearized Large Language Models}, author = {Haoran You and Yichao Fu and Zheng Wang and Amir Yazdanbakhsh and Yingyan Celine Lin}, year = {2024}, url = {https://arxiv.org/abs/2406.07368}, note = {Source identifier: 2406.07368} }