- port retrieval, validation, and eval improvements relevant to os - align prompts and dimensions with the flat single-agent model - replace the old eval suite with the focused core scenarios Generated with Codex
14 lines
450 B
JSON
14 lines
450 B
JSON
{
|
|
"id": "golden-ra-h-core-v2",
|
|
"name": "Golden RA-H Core v2",
|
|
"description": "Slim 5-scenario eval set for core RA-H usage: focused graph writes, skill-guided writes, indexed node search, chunk-grounded insight creation, and hub traversal.",
|
|
"inherits": "golden-v1",
|
|
"version": 2,
|
|
"focus": [
|
|
"tool-call correctness",
|
|
"skill usage only when appropriate",
|
|
"core retrieval/write behavior",
|
|
"latency/tokens/cost guards"
|
|
]
|
|
}
|