@misc{indiciae422bb2468256, title = {LAVCap: LLM-based Audio-Visual Captioning using Optimal Transport}, author = {Kyeongha Rho and Hyeongkeun Lee and Valentio Iverson and Joon Son Chung}, year = {2025}, url = {https://arxiv.org/abs/2501.09291}, note = {Source identifier: 2501.09291} }