@misc{indiciae2ec13c8f3442, title = {BLIVA: A Simple Multimodal LLM for Better Handling of Text-Rich Visual Questions}, author = {Wenbo Hu and Yifan Xu and Yi Li and Weiyue Li and Zeyuan Chen and Zhuowen Tu}, year = {2023}, url = {https://arxiv.org/abs/2308.09936}, note = {Source identifier: 2308.09936} }