@misc{indiciae1bfaf579a72e, title = {EventLens: Leveraging Event-Aware Pretraining and Cross-modal Linking Enhances Visual Commonsense Reasoning}, author = {Mingjie Ma and Zhihuan Yu and Yichao Ma and Guohui Li}, year = {2024}, url = {https://arxiv.org/abs/2404.13847}, note = {Source identifier: 2404.13847} }