@misc{indiciae7d85a655ef2a, title = {Efficient Prompt Compression with Evaluator Heads for Long-Context Transformer Inference}, author = {Weizhi Fei and Xueyan Niu and Guoqing Xie and Yingqing Liu and Bo Bai and Wei Han}, year = {2025}, url = {https://arxiv.org/abs/2501.12959}, note = {Source identifier: 2501.12959} }