@misc{indiciae9a512bd249d3, title = {Macaw-LLM: Multi-Modal Language Modeling with Image, Audio, Video, and Text Integration}, author = {Chenyang Lyu and Minghao Wu and Longyue Wang and Xinting Huang and Bingshuai Liu and Zefeng Du and Shuming Shi and Zhaopeng Tu}, year = {2023}, url = {https://arxiv.org/abs/2306.09093}, note = {Source identifier: 2306.09093} }