@misc{indiciae990702dc51b5, title = {Align before Attend: Aligning Visual and Textual Features for Multimodal Hateful Content Detection}, author = {Eftekhar Hossain and Omar Sharif and Mohammed Moshiul Hoque and Sarah M. Preum}, year = {2024}, url = {https://arxiv.org/abs/2402.09738}, note = {Source identifier: 2402.09738} }