@misc{indiciaeba5616e9b930, title = {Towards Multimodal Vision-Language Models Generating Non-Generic Text}, author = {Wes Robbins and Zanyar Zohourianshahzadi and Jugal Kalita}, year = {2022}, url = {https://arxiv.org/abs/2207.04174}, note = {Source identifier: 2207.04174} }