@misc{indiciae82b0d21c289d, title = {Speed and Conversational Large Language Models: Not All Is About Tokens per Second}, author = {Javier Conde and Miguel González and Pedro Reviriego and Zhen Gao and Shanshan Liu and Fabrizio Lombardi}, year = {2025}, doi = {10.1109/mc.2024.3399384}, url = {https://arxiv.org/abs/2502.16721}, note = {Source identifier: 2502.16721} }