@misc{indiciaebc3e7e577014, title = {Stacked Cross-modal Feature Consolidation Attention Networks for Image Captioning}, author = {Mozhgan Pourkeshavarz and Shahabedin Nabavi and Mohsen Ebrahimi Moghaddam and Mehrnoush Shamsfard}, year = {2023}, doi = {10.1007/s11042-023-15869-x}, url = {https://arxiv.org/abs/2302.04676}, note = {Source identifier: 2302.04676} }