@misc{indiciae1e14bf77f2fb, title = {Enhancing and Accelerating Large Language Models via Instruction-Aware Contextual Compression}, author = {Haowen Hou and Fei Ma and Binwen Bai and Xinxin Zhu and Fei Yu}, year = {2024}, url = {https://arxiv.org/abs/2408.15491}, note = {Source identifier: 2408.15491} }