@misc{indiciae11777ff90832, title = {Efficient Large Multi-modal Models via Visual Context Compression}, author = {Jieneng Chen and Luoxin Ye and Ju He and Zhao-Yang Wang and Daniel Khashabi and Alan Yuille}, year = {2024}, url = {https://arxiv.org/abs/2406.20092}, note = {Source identifier: 2406.20092} }