Glenn Matlin, Isaac Song, Anthony Wen-Ming Zang, and Mark O. Riedl No One Wins in Nuclear War: Social Simulations of High-Stakes Military Decision-Making Proceedings of the COLM 2026 Workshop on Social Simulations with LLMs (2026). arXivWorkshopbibtexSocial SimulationAgentsLarge Language ModelsAI Safety
@InProceedings{Matlin2026NoOneWins,
author = {Matlin, Glenn and Song, Isaac and Zang, Anthony Wen-Ming and Riedl, Mark O.},
booktitle = {Proceedings of the COLM 2026 Workshop on Social Simulations with LLMs},
title = {No One Wins in Nuclear War: Social Simulations of High-Stakes Military Decision-Making},
year = {2026},
owner = {riedl},
url = {https://arxiv.org/abs/2608.01868},
keywords = {worlds, agents, llms, safe},
}
Mark O. Riedl and Glenn Matlin Position: AI is Not Ready for Strategic Conflict Proceedings of the COLM 2026 Workshop on Social Simulations with LLMs (2026). WorkshopbibtexSocial SimulationAgentsPolicy and LawAI Safety
@InProceedings{Riedl2026Position,
author = {Riedl, Mark O. and Matlin, Glenn},
booktitle = {Proceedings of the COLM 2026 Workshop on Social Simulations with LLMs},
title = {Position: AI is Not Ready for Strategic Conflict},
year = {2026},
owner = {riedl},
keywords = {worlds, agents, policy, safe},
}
Geigh Zollicoffer, Tanush Chopra, Mingkuan Yan, Xiaoxu Ma, Kenneth Eaton, and Mark Riedl World Model Robustness via Surprise Recognition Proceedings of the 2026 CVPR Findings (2026). arXivConferencebibtexAgentsReinforcement LearningAI Safety
@InProceedings{Zollicoffer2026WorldModel,
author = {Zollicoffer, Geigh and Tanush Chopra and Mingkuan Yan and Xiaoxu Ma and Kenneth Eaton and Mark Riedl},
booktitle = {Proceedings of the 2026 CVPR Findings},
title = {World Model Robustness via Surprise Recognition},
year = {2026},
owner = {riedl},
url = {https://arxiv.org/abs/2512.01119},
keywords = {agents, rl, safe},
}
2025
Glenn Matlin, Parv Mahajan, Isaac Song, Yixiong Hao, Ryan Bard, Stu Topp, Evan Montoya, M. Rehan Parwani, Soham Shetty, and Mark Riedl Shall We Play a Game? Language Models for Open-ended Wargames Proceedings of the 2025 EMNLP Word Play Workshop (2025). arXivWorkshopbibtexSocial SimulationAgentsAI Safety
@InProceedings{Matlin2025Wargames,
author = {Glenn Matlin and Parv Mahajan and Isaac Song and Yixiong Hao and Ryan Bard and Stu Topp and Evan Montoya and Parwani, M. Rehan and Soham Shetty and Mark Riedl},
booktitle = {Proceedings of the 2025 EMNLP Word Play Workshop},
title = {Shall We Play a Game? Language Models for Open-ended Wargames},
year = {2025},
owner = {riedl},
url = {https://arxiv.org/abs/2509.17192},
keywords = {worlds, agents, safe},
}
Geigh Zollicoffer, Kenneth Eaton, Jonathan Balloch, Julia Kim, Riedl Mark O., and Robert Wright Novelty Detection in Reinforcement Learning with World Models International Conference on Machine Learning (2025). arXivConferencebibtexAgentsReinforcement LearningAI Safety
@InProceedings{Zollicoffer2023Novelty,
author = {Geigh Zollicoffer and Kenneth Eaton and Jonathan Balloch and Julia Kim and Riedl Mark O. and Robert Wright},
booktitle = {International Conference on Machine Learning},
title = {Novelty Detection in Reinforcement Learning with World Models},
year = {2025},
owner = {riedl},
url = {https://arxiv.org/abs/2310.08731},
keywords = {agents, rl, safe},
}
@InProceedings{Baheti2023LOL,
author = {Ashutosh Baheti and Ximing Lu and Faeze Brahman and Ronan Le Bras and Maarten Sap and Riedl, Mark O.},
booktitle = {Proceedings of ICLR 2024},
title = {Improving Language Models with Advantage-based Offline Policy Gradients},
year = {2024},
owner = {riedl},
url = {https://arxiv.org/abs/2305.14718},
keywords = {llms, rl, align, safe},
}
2021
Ashutosh Baheti, Maarten Sap, Alan Ritter, and Mark O. Riedl Just Say No: Analyzing the Stance of Neural Dialogue Generation in Offensive Contexts Proceedings of EMNLP 2021 (2021). arXivConferencebibtexLarge Language ModelsValue AlignmentAI Safety
@InProceedings{Baheti2021JustSayNo,
author = {Ashutosh Baheti and Maarten Sap and Alan Ritter and Riedl, Mark O.},
title = {Just Say No: Analyzing the Stance of Neural Dialogue Generation in Offensive Contexts},
booktitle = {Proceedings of EMNLP 2021},
year = {2021},
owner = {riedl},
timestamp = {2021.09.01},
url = {https://arxiv.org/abs/2108.11830},
keywords = {llms, align, safe},
}
@Article{NahianTraining2021,
author = {Md Sultan Al Nahian and Spencer Frazier and Brent Harrison and Riedl, Mark O.},
title = {Training Value-Aligned Reinforcement Learning Agents Using a Normative Prior},
journal = {arXiv:2104.09469},
year = {2021},
owner = {riedl},
timestamp = {2021.05.21},
url = {https://arxiv.org/abs/2104.09469},
keywords = {agents, rl, align, safe},
}
2020
Xiangyu Peng, S. Li, Spencer Frazier, and Mark O. Riedl Reducing Non-Normative Text Generation from Language Models International Conference on Natural Language Generation (2020). arXivConferencebibtexLarge Language ModelsValue AlignmentAI Safety
@InProceedings{Peng2020ReducingNT,
author = {Xiangyu Peng and S. Li and Spencer Frazier and Mark O. Riedl},
title = {Reducing Non-Normative Text Generation from Language Models},
booktitle = {International Conference on Natural Language Generation},
year = {2020},
url = {https://arxiv.org/abs/2001.08764},
keywords = {llms, align, safe},
}
2017
Mark O. Riedl and Brent. Harrison Enter the Matrix: A Virtual World Approach to Safely Interruptable Autonomous Systems Proceedings of the AAAI 2017 Workshop on SafeAI (2017). arXivWorkshopbibtexAgentsReinforcement LearningAI Safety
@InProceedings{riedl:arxiv:matrix2017,
author = {Riedl, Mark~O. and Harrison, Brent.},
title = {{Enter the Matrix: A Virtual World Approach to Safely Interruptable Autonomous Systems}},
booktitle = {Proceedings of the AAAI 2017 Workshop on SafeAI},
year = {2017},
url = {https://arxiv.org/abs/1703.10284},
keywords = {agents, rl, safe},
}
2016
Mark Riedl and Brent Harrison Using Stories to Teach Human Values to Artificial Agents Proceedings of the 2nd International Workshop on AI, Ethics and Society (2016). PDFWorkshopbibtexAgentsValue AlignmentAI Safety
@InProceedings{riedl:aaai-ethics2016,
author = {Riedl, Mark and Harrison, Brent},
title = {Using Stories to Teach Human Values to Artificial Agents},
booktitle = {Proceedings of the 2nd International Workshop on {AI}, Ethics and Society},
year = {2016},
owner = {riedl},
timestamp = {2015.10.23},
url = {http://www.cc.gatech.edu/~riedl/pubs/aaai-ethics16.pdf},
keywords = {agents, align, safe},
}
Brent Harrsion and Mark O Riedl Learning From Stories: Using Crowdsourced Narratives to Train Virtual Agents Proceedings of the 2016 AAAI Conference on Artificial Intelligence and Interactive Digital Entertainment (2016). PDFConferencebibtexAgentsValue AlignmentAI SafetyHuman Computation
@InProceedings{harrison:aiide2016,
author = {Harrsion, Brent and Riedl, Mark O},
title = {Learning From Stories: Using Crowdsourced Narratives to Train Virtual Agents},
booktitle = {Proceedings of the 2016 AAAI Conference on Artificial Intelligence and Interactive Digital Entertainment},
year = {2016},
owner = {riedl},
timestamp = {2016.06.08},
url = {http://www.cc.gatech.edu/~riedl/pubs/harrison-aiide16.pdf},
keywords = {agents, align, safe, hcomp},
}
Brent Harrison and Mark Riedl Towards Learning From Stories: An Approach for Interactive Machine Learning Proceedings of the AAAI'16 Workshop on Symbiotic Cognitive Systems (2016). PDFWorkshopbibtexAgentsValue AlignmentAI Safety
@InProceedings{harrison:aaai-symbiotic2016,
author = {Harrison, Brent and Riedl, Mark},
title = {Towards Learning From Stories: An Approach for Interactive Machine Learning},
booktitle = {Proceedings of the AAAI'16 Workshop on Symbiotic Cognitive Systems},
year = {2016},
owner = {riedl},
timestamp = {2015.10.23},
url = {http://www.cc.gatech.edu/~riedl/pubs/aaai-symbiotic16.pdf},
keywords = {agents, align, safe},
}
Brent Harrison, Siddhartha Banerjee, and Mark O. Riedl Learning from Stories: Using Natural Communication to Train Believable Agents Proceedings of the 2016 IJCAI Workshop on Interactive Machine Learning (2016). PDFWorkshopbibtexAgentsValue AlignmentAI Safety
@InProceedings{harrison:ijcai-iml2016,
author = {Harrison, Brent and Banerjee, Siddhartha and Riedl, Mark O.},
title = {Learning from Stories: Using Natural Communication to Train Believable Agents},
booktitle = {Proceedings of the 2016 IJCAI Workshop on Interactive Machine Learning},
year = {2016},
owner = {riedl},
timestamp = {2015.11.21},
url = {http://www.cc.gatech.edu/~riedl/pubs/ijcai-iml16.pdf},
keywords = {agents, align, safe},
}