{"kind":"task","effective_mode":"full","benchmark":{"kind":"benchmark","effective_mode":"full","slug":"cruxeval","formal_name":"CRUXEval","introduction":"CRUXEval takes 800 short Python functions and asks either for the output given an input (CRUXEval-O) or for an input that produces a given output (CRUXEval-I). It measures the ability to execute code mentally rather than to write it.","introduction_ja":"","introduction_en":"","category":"Category not supplied","task_count":null,"acquisition_status":"Acquisition status not supplied","official_url":"https://crux-eval.github.io/","indexing_mode":"noindex","profile":{"resources":[],"task_format":"","scoring":"","metric":"","size":"","answer_access":"","license":"","citation":"","maintainer":"","released":"","why_hard":"","related":[]}},"task_id":"f1606e07-b023-5c45-9f9b-dfac550670c6","task_key":"default--test--sample~5f111","task_revision_id":"2","upstream_id":"sample_111","short_description":"{'x': 67, 'v': 89, '': 4, 'alij': 11, 'kgfsd': 72, 'yafby': 83}","config":"default","split":"test","body":"{\"code\":\"def f(marks):\\n    highest = 0\\n    lowest = 100\\n    for value in marks.values():\\n        if value > highest:\\n            highest = value\\n        if value < lowest:\\n            lowest = value\\n    return highest, lowest\",\"input\":\"{'x': 67, 'v': 89, '': 4, 'alij': 11, 'kgfsd': 72, 'yafby': 83}\"}","display_format":"code","language":"","answer_status":"published","assets":[],"source_url":"https://crux-eval.github.io/","history":"initial import","indexing_mode":"noindex","subproblems":[],"grids":[]}