@misc{indiciae6a99916e178e, title = {Effectively Enhancing Vision Language Large Models by Prompt Augmentation and Caption Utilization}, author = {Minyi Zhao and Jie Wang and Zhaoyang Li and Jiyuan Zhang and Zhenbang Sun and Shuigeng Zhou}, year = {2024}, url = {https://arxiv.org/abs/2409.14484}, note = {Source identifier: 2409.14484} }