This dataset contains the single-choice (mc1) task from the original TruthfulQA benchmark.
The answer options were shuffled, so the correct answer does not always appear first.
@misc{lin2021truthfulqa,
title={TruthfulQA: Measuring How Models Mimic Human Falsehoods},
author={Stephanie Lin and Jacob Hilton and Owain Evans},
year={2021},
eprint={2109.07958},
archivePrefix={arXiv},
primaryClass={cs.CL}
}