Publications
2026
-
ICMLIntrinsic Credit Assignment for Long Horizon InteractionInternational Conference on Machine Learning (ICML), 2026
@inproceedings{auzina2026intrinsic, title={Intrinsic Credit Assignment for Long Horizon Interaction}, author={Auzina, Ilze Amanda and Str{\"u}ber, Joschka and Hern{\'a}ndez-Guti{\'e}rrez, Sergio and Goel, Shashwat and Prabhu, Ameya and Bethge, Matthias}, booktitle={Proceedings of the 43rd International Conference on Machine Learning}, year={2026} } -
ICML-WRevengeBench: Reverse Engineering Code-Space Policies from Behavioral ExperimentsICML 2026 Workshop on Agents in the Wild: Safety, Security, and Beyond (AIWILD)
@inproceedings{rahmani2026revengebench, title={RevengeBench: Reverse Engineering Code-Space Policies from Behavioral Experiments}, author={Rahmani, Babak and Dziadzio, Sebastian and Str{\"u}ber, Joschka and Hern{\'a}ndez-Guti{\'e}rrez, Sergio and Bethge, Matthias}, booktitle={ICML 2026 Workshop on Agents in the Wild: Safety, Security, and Beyond (AIWILD)}, year={2026} } -
ICML-WQVal: Cheaply Evaluating Dense Supervision Signals for Long-Horizon LLM AgentsICML 2026 Workshop on Reinforcement Learning from World Feedback (RLxF)
@inproceedings{hernandezgutierrez2026qval, title={QVal: Cheaply Evaluating Dense Supervision Signals for Long-Horizon LLM Agents}, author={Hern{\'a}ndez-Guti{\'e}rrez, Sergio and Merler, Matteo and Auzina, Ilze Amanda and Str{\"u}ber, Joschka and Prabhu, Ameya and Bethge, Matthias}, booktitle={ICML 2026 Workshop on Reinforcement Learning from World Feedback (RLxF)}, year={2026} } -
arXivDataComp-VLM: Improved Open Datasets for Vision-Language ModelsarXiv Preprint
@misc{farina2026datacompvlm, title={DataComp-VLM: Improved Open Datasets for Vision-Language Models}, author={Farina, Matteo and Udandarao, Vishaal and Nguyen, Thao and Kuzucu, Selim and B{\"o}ther, Maximilian and Hochlehnert, Andreas and Ghosh, Adhiraj and Nezhurina, Marianna and Roth, Karsten and Str{\"u}ber, Joschka and Zhang, Yuhui and Dziadzio, Sebastian and Sui, Elaine and Jahagirdar, Soumya and Ghosh, Dhruba and Hammoud, Hasan and De Min, Thomas and Caldarella, Simone and Mirza, Jehanzeb and Keh, Sedrick and Cherti, Mehdi and Kuehne, Hilde and Schiele, Bernt and Yeung-Levy, Serena and Naeem, Muhammad Ferjad and Tombari, Federico and Klimovic, Ana and Ricci, Elisa and Bethge, Matthias and Oh, Sewoong and Prabhu, Ameya and Tonioni, Alessio and Jitsev, Jenia and Mancini, Massimiliano and Schmidt, Ludwig and Parthasarathy, Nikhil}, year={2026}, eprint={2606.28551}, archivePrefix={arXiv}, primaryClass={cs.CV}, url={https://arxiv.org/abs/2606.28551} }
2025
-
ICML-W OralMeasuring Belief Updates in Curious AgentsICML 2025 Workshop on Assessing World Models
@inproceedings{strueber2025measuring, title={Measuring Belief Updates in Curious Agents}, author={Str{\"u}ber, Joschka and Auzina, Ilze Amanda and Goel, Shashwat and Keller, Susanne and Geiping, Jonas and Prabhu, Ameya and Bethge, Matthias}, booktitle={ICML 2025 Workshop on Assessing World Models}, year={2025} } -
ICML SpotlightGreat Models Think Alike and this Undermines AI OversightInternational Conference on Machine Learning (ICML), 2025
@inproceedings{goel2025great, title={Great Models Think Alike and this Undermines AI Oversight}, author={Goel, Shashwat and Str{\"u}ber, Joschka and Auzina, Ilze Amanda and Chandra, Karuna K and Kumaraguru, Ponnurangam and Kiela, Douwe and Prabhu, Ameya and Bethge, Matthias and Geiping, Jonas}, booktitle={Proceedings of the 42nd International Conference on Machine Learning}, year={2025} }