mirror of
https://github.com/open-compass/opencompass.git
synced 2025-05-30 16:03:24 +08:00

* bigcodebench * humaneval * humanevalx * humanevalx * livecodebench * mbpp * humaneval_plus * fix bug * template * max_out fix * template update
44 lines
1.5 KiB
Python
44 lines
1.5 KiB
Python
from opencompass.openicl.icl_prompt_template import PromptTemplate
|
|
from opencompass.openicl.icl_retriever import ZeroRetriever
|
|
from opencompass.openicl.icl_inferencer import GenInferencer
|
|
from opencompass.datasets import HumanevalXDataset, HumanevalXEvaluator
|
|
|
|
humanevalx_reader_cfg = dict(
|
|
input_columns=['prompt'], output_column='declaration', train_split='test')
|
|
|
|
humanevalx_infer_cfg = dict(
|
|
prompt_template=dict(
|
|
type=PromptTemplate,
|
|
template='{prompt}'),
|
|
retriever=dict(type=ZeroRetriever),
|
|
inferencer=dict(type=GenInferencer))
|
|
|
|
humanevalx_eval_cfg_dict = {
|
|
lang : dict(
|
|
evaluator=dict(
|
|
type=HumanevalXEvaluator,
|
|
language=lang,
|
|
ip_address=
|
|
'localhost', # replace to your code_eval_server ip_address, port
|
|
port=5001), # refer to https://opencompass.readthedocs.io/en/latest/advanced_guides/code_eval_service.html to launch a server
|
|
pred_role='BOT')
|
|
for lang in ['python', 'cpp', 'go', 'java', 'js'] # do not support rust now
|
|
}
|
|
|
|
# Please download the needed `xx.jsonl.gz` from
|
|
# https://github.com/THUDM/CodeGeeX2/tree/main/benchmark/humanevalx
|
|
# and move them into `data/humanevalx/` folder
|
|
humanevalx_datasets = [
|
|
dict(
|
|
type=HumanevalXDataset,
|
|
abbr=f'humanevalx-{lang}',
|
|
language=lang,
|
|
path='./data/humanevalx',
|
|
reader_cfg=humanevalx_reader_cfg,
|
|
infer_cfg=humanevalx_infer_cfg,
|
|
eval_cfg=humanevalx_eval_cfg_dict[lang],
|
|
n=5,
|
|
k=3)
|
|
for lang in ['python', 'cpp', 'go', 'java', 'js']
|
|
]
|