@misc{indiciae784aca99d6cf, title = {ViCLEVR: A Visual Reasoning Dataset and Hybrid Multimodal Fusion Model for Visual Question Answering in Vietnamese}, author = {Khiem Vinh Tran and Hao Phu Phan and Kiet Van Nguyen and Ngan Luu Thuy Nguyen}, year = {2023}, url = {https://arxiv.org/abs/2310.18046}, note = {Source identifier: 2310.18046} }