@misc{indiciae47368e2ca0a0, title = {SnapNTell: Enhancing Entity-Centric Visual Question Answering with Retrieval Augmented Multimodal LLM}, author = {Jielin Qiu and Andrea Madotto and Zhaojiang Lin and Paul A. Crook and Yifan Ethan Xu and Xin Luna Dong and Christos Faloutsos and Lei Li and Babak Damavandi and Seungwhan Moon}, year = {2024}, url = {https://arxiv.org/abs/2403.04735}, note = {Source identifier: 2403.04735} }