@misc{indiciaef400324ce774, title = {Linear Alignment of Vision-language Models for Image Captioning}, author = {Fabian Paischer and Markus Hofmarcher and Sepp Hochreiter and Thomas Adler}, year = {2025}, url = {https://arxiv.org/abs/2307.05591}, note = {Source identifier: 2307.05591} }