Published in: LIPIcs, Volume 394, 20th International Symposium on Empirical Software Engineering and Measurement (ESEM 2026)
Hui Sun, Anderson Uchôa, Rohit Gheyi, and Wesley K. G. Assunção. Evaluating Language Models on Cross-Language Code Functional Equivalence. In 20th International Symposium on Empirical Software Engineering and Measurement (ESEM 2026). Leibniz International Proceedings in Informatics (LIPIcs), Volume 394, pp. 41:1-41:21, Schloss Dagstuhl – Leibniz-Zentrum für Informatik (2026)
@InProceedings{sun_et_al:LIPIcs.ESEM.2026.41,
author = {Sun, Hui and Uch\^{o}a, Anderson and Gheyi, Rohit and Assun\c{c}\~{a}o, Wesley K. G.},
title = {{Evaluating Language Models on Cross-Language Code Functional Equivalence}},
booktitle = {20th International Symposium on Empirical Software Engineering and Measurement (ESEM 2026)},
pages = {41:1--41:21},
series = {Leibniz International Proceedings in Informatics (LIPIcs)},
ISBN = {978-3-95977-450-5},
ISSN = {1868-8969},
year = {2026},
volume = {394},
editor = {Feldt, Robert and Paasivaara, Maria and Mendez, Daniel and Wagner, Stefan and Bar\'{o}n, Marvin Mu\~{n}oz},
publisher = {Schloss Dagstuhl -- Leibniz-Zentrum f{\"u}r Informatik},
address = {Dagstuhl, Germany},
URL = {https://drops.dagstuhl.de/entities/document/10.4230/LIPIcs.ESEM.2026.41},
URN = {urn:nbn:de:0030-drops-280091},
doi = {10.4230/LIPIcs.ESEM.2026.41},
annote = {Keywords: Empirical software engineering, code functional equivalence, generative AI, cross-language analysis, benchmark, failure analysis}
}