@misc{indiciaeebeb922bab3f, title = {AirCache: Activating Inter-modal Relevancy KV Cache Compression for Efficient Large Vision-Language Model Inference}, author = {Kai Huang and Hao Zou and Bochen Wang and Ye Xi and Zhen Xie and Hao Wang}, year = {2025}, url = {https://arxiv.org/abs/2503.23956}, note = {Source identifier: 2503.23956} }