@misc{indiciae9946a111e40d, title = {Optimizing Vision-Language Interactions Through Decoder-Only Models}, author = {Kaito Tanaka and Benjamin Tan and Brian Wong}, year = {2024}, url = {https://arxiv.org/abs/2412.10758}, note = {Source identifier: 2412.10758} }