hermes-agent/tests/tools/test_schema_sanitizer.py
Teknium 39975613b1
test: prune wave 2 + speed fixes — 28,106 → 19,757 test functions, suite wall 315s → 294s
Second, deeper pass over tools/gateway/hermes_cli plus first pass over
the trees wave 1 missed (acp, acp_adapter, skills, computer_use, docker,
dashboard, conformance, monitoring, secret_sources, hermes_state,
providers). Same rubric as wave 1 (AGENTS.md test policy); security,
alternation/caching invariants, issue-number regressions, and E2E kept.

Real test-quality fixes found and rooted out along the way:
- tests/tools/test_command_guards.py made real auxiliary-LLM HTTPS calls
  (DEFAULT_CONFIG smart-approval leaked in) — pinned approval
  mode=manual via autouse fixture: 17.4s → 0.4s.
- test_model_switch_custom_providers.py / test_user_providers_model_switch.py
  silently probed live provider catalogs (~2s/test) — stubbed
  cached_provider_model_ids/provider_model_ids/fetch_api_models.
- test_telegram_noise_filter.py: 15-platform copy-paste matrix over
  shared gateway.run logic → 3 representative platforms (55s → 3.9s).
- test_gateway_shutdown.py: stop()'s 5s interrupt-deadline loop spun on
  MagicMock agents — interrupt.side_effect now clears _running_agents
  (22s → 1.0s).
- test_gateway_inactivity_timeout.py poll-harness timings shrunk 3-5x
  (24s → 1.1s); test_mcp_stability.py backoff/SIGTERM-grace sleeps
  patched (15.4s → 2.5s); test_async_delegation.py negative-drain wait
  5s → 0.5s.
- test_telegram_init_deadline.py: loop-block margin restored to 1.0s
  with rationale comment — the watchdog-dump assertion needs the loop
  blocked well past deadline+grace under parallel load (flaked once in
  the 40-worker verification run at a 0.2s margin).

Verification: full hermetic suite via scripts/run_tests.sh —
2,438 files, 21,718 tests passed, 0 failed, 293.9s wall.
Suite totals vs original baseline: 46,820 → 19,757 test functions
(−57.8%), wall 583.5s → 293.9s (−50%), subprocess CPU 13,564s → 11,623s.
2026-07-29 13:39:40 -07:00

348 lines
13 KiB
Python

"""Tests for tools/schema_sanitizer.py.
Targets the known llama.cpp ``json-schema-to-grammar`` failure modes that
cause ``HTTP 400: Unable to generate parser for this template. ...
Unrecognized schema: "object"`` errors on local inference backends.
"""
from __future__ import annotations
import copy
from tools.schema_sanitizer import (
sanitize_tool_schemas,
strip_pattern_and_format,
strip_slash_enum,
)
def _tool(name: str, parameters: dict) -> dict:
return {"type": "function", "function": {"name": name, "parameters": parameters}}
def test_object_without_properties_gets_empty_properties():
tools = [_tool("t", {"type": "object"})]
out = sanitize_tool_schemas(tools)
assert out[0]["function"]["parameters"] == {"type": "object", "properties": {}}
def test_nested_object_without_properties_gets_empty_properties():
tools = [_tool("t", {
"type": "object",
"properties": {
"name": {"type": "string"},
"arguments": {"type": "object", "description": "free-form"},
},
"required": ["name"],
})]
out = sanitize_tool_schemas(tools)
args = out[0]["function"]["parameters"]["properties"]["arguments"]
assert args["type"] == "object"
assert args["properties"] == {}
assert args["description"] == "free-form"
def test_bare_string_object_value_replaced_with_schema_dict():
# Malformed: a property's schema value is the bare string "object".
# This is the exact shape llama.cpp reports as `Unrecognized schema: "object"`.
tools = [_tool("t", {
"type": "object",
"properties": {
"payload": "object", # <-- invalid, should be {"type": "object"}
},
})]
out = sanitize_tool_schemas(tools)
payload = out[0]["function"]["parameters"]["properties"]["payload"]
assert isinstance(payload, dict)
assert payload["type"] == "object"
assert payload["properties"] == {}
def test_nullable_type_array_collapsed_to_single_string():
tools = [_tool("t", {
"type": "object",
"properties": {
"maybe_name": {"type": ["string", "null"]},
},
})]
out = sanitize_tool_schemas(tools)
prop = out[0]["function"]["parameters"]["properties"]["maybe_name"]
assert prop["type"] == "string"
assert prop.get("nullable") is True
def test_multitype_array_becomes_anyof_no_branch_dropped():
# Ported from anomalyco/opencode#31877: a genuine multi-type array such as
# ["number", "string"] (common in MCP tool schemas) must keep BOTH branches
# as an anyOf, not silently drop all but the first. Several backends
# (llama.cpp, Gemini via OpenAI-compatible transports) reject the array form.
tools = [_tool("t", {
"type": "object",
"properties": {
"status": {"type": ["number", "string"], "description": "status filter"},
},
})]
out = sanitize_tool_schemas(tools)
prop = out[0]["function"]["parameters"]["properties"]["status"]
assert "type" not in prop
assert prop["anyOf"] == [{"type": "number"}, {"type": "string"}]
assert prop.get("nullable") is None
# Sibling keywords survive alongside the generated anyOf.
assert prop["description"] == "status filter"
def test_all_null_type_array_becomes_null_type():
tools = [_tool("t", {
"type": "object",
"properties": {
"n": {"type": ["null"]},
},
})]
out = sanitize_tool_schemas(tools)
prop = out[0]["function"]["parameters"]["properties"]["n"]
assert prop["type"] == "null"
def test_single_element_type_array_unwrapped():
tools = [_tool("t", {
"type": "object",
"properties": {
"s": {"type": ["string"]},
},
})]
out = sanitize_tool_schemas(tools)
prop = out[0]["function"]["parameters"]["properties"]["s"]
assert prop["type"] == "string"
assert prop.get("nullable") is None
def test_anyof_nested_objects_sanitized():
tools = [_tool("t", {
"type": "object",
"properties": {
"opt": {
"anyOf": [
{"type": "object"}, # bare object
{"type": "string"},
],
},
},
})]
out = sanitize_tool_schemas(tools)
variants = out[0]["function"]["parameters"]["properties"]["opt"]["anyOf"]
assert variants[0] == {"type": "object", "properties": {}}
assert variants[1] == {"type": "string"}
def test_missing_parameters_gets_default_object_schema():
tools = [{"type": "function", "function": {"name": "t"}}]
out = sanitize_tool_schemas(tools)
assert out[0]["function"]["parameters"] == {"type": "object", "properties": {}}
def test_non_dict_parameters_gets_default_object_schema():
tools = [_tool("t", "object")] # pathological
out = sanitize_tool_schemas(tools)
assert out[0]["function"]["parameters"] == {"type": "object", "properties": {}}
def test_required_pruned_to_existing_properties():
tools = [_tool("t", {
"type": "object",
"properties": {"name": {"type": "string"}},
"required": ["name", "missing_field"],
})]
out = sanitize_tool_schemas(tools)
assert out[0]["function"]["parameters"]["required"] == ["name"]
def test_well_formed_schema_unchanged():
schema = {
"type": "object",
"properties": {
"path": {"type": "string", "description": "File path"},
"offset": {"type": "integer", "minimum": 1},
},
"required": ["path"],
}
tools = [_tool("read_file", copy.deepcopy(schema))]
out = sanitize_tool_schemas(tools)
assert out[0]["function"]["parameters"] == schema
def test_additional_properties_schema_sanitized():
tools = [_tool("t", {
"type": "object",
"properties": {
"dict_field": {
"type": "object",
"additionalProperties": {"type": "object"}, # bare object schema
},
},
})]
out = sanitize_tool_schemas(tools)
field = out[0]["function"]["parameters"]["properties"]["dict_field"]
assert field["additionalProperties"] == {"type": "object", "properties": {}}
def test_items_sanitized_in_array_schema():
tools = [_tool("t", {
"type": "object",
"properties": {
"bag": {
"type": "array",
"items": {"type": "object"}, # bare object items
},
},
})]
out = sanitize_tool_schemas(tools)
items = out[0]["function"]["parameters"]["properties"]["bag"]["items"]
assert items == {"type": "object", "properties": {}}
# ─────────────────────────────────────────────────────────────────────────
# strip_pattern_and_format — reactive recovery when llama.cpp rejects a
# schema with an HTTP 400 grammar-parse error. Must be opt-in (only
# invoked on recovery) and must not damage property names.
# ─────────────────────────────────────────────────────────────────────────
def test_strip_responses_mixed_formats():
"""Mixed list of OpenAI-format and Responses-format tools should both be sanitized."""
from tools.schema_sanitizer import strip_pattern_and_format
tools = [
# OpenAI-format: {"function": {"parameters": {...}}}
{
"type": "function",
"function": {
"name": "search",
"parameters": {
"type": "object",
"properties": {
"query": {"type": "string", "pattern": "^[a-z]+$"}
}
}
}
},
# Responses-format: {"name": "...", "parameters": {...}}
{
"name": "get_time",
"parameters": {
"type": "object",
"properties": {
"tz": {"type": "string", "format": "date-time"}
}
},
"type": "function"
}
]
result, stripped = strip_pattern_and_format(tools)
assert stripped == 2, f"Expected 2 stripped (1 pattern + 1 format), got {stripped}"
# OpenAI-format tool: pattern stripped from parameters
openai_params = result[0]["function"]["parameters"]["properties"]["query"]
assert "pattern" not in openai_params, f"pattern should be stripped: {openai_params}"
# Responses-format tool: format stripped
resp_params = result[1]["parameters"]["properties"]["tz"]
assert "format" not in resp_params, f"format should be stripped: {resp_params}"
# Verify structure preserved
assert result[0]["function"]["parameters"]["type"] == "object"
assert result[1]["parameters"]["type"] == "object"
# ─────────────────────────────────────────────────────────────────────────
# strip_slash_enum — reactive recovery when xAI's /v1/responses (and
# /v1/chat/completions) grammar-compiler rejects enum values containing
# a forward slash. Symptom: HTTP 400 "Invalid arguments passed to the
# model" before any token is emitted. Most commonly hit by MCP-derived
# tools whose enum lists HuggingFace IDs like "Qwen/Qwen3.5-0.8B".
# ─────────────────────────────────────────────────────────────────────────
# ---------------------------------------------------------------------------
# Property-key renaming (provider ^[a-zA-Z0-9_.-]{1,64}$ pattern compat)
# Real-world source: Cloudflare flat API MCP ships keys like
# ``issue_class~neq`` and ``meta.<field>[<operator>]`` — one bad key anywhere
# in the tools array 400s the whole request on Anthropic/Bedrock/Vertex/Azure.
# ---------------------------------------------------------------------------
from tools.schema_sanitizer import sanitize_property_key, unrename_tool_args
def test_sanitize_property_key_empty_falls_back():
assert sanitize_property_key("~~~") == "___"
assert sanitize_property_key("") == "param"
# ---------------------------------------------------------------------------
# dependentRequired -- literal property-name strings must survive
# ---------------------------------------------------------------------------
def test_dependent_required_preserved_through_public_api():
"""dependentRequired values are literal property names, not schemas."""
schema = {
"type": "object",
"properties": {
"owner": {"type": "string"},
"repo": {"type": "string"},
"organization": {"type": "string"},
},
"dependentRequired": {
"owner": ["repo", "organization"],
"repo": ["owner"],
},
}
tools = [_tool("t", copy.deepcopy(schema))]
out = sanitize_tool_schemas(tools)
params = out[0]["function"]["parameters"]
dep = params.get("dependentRequired", {})
# Values are the original property-name strings unchanged.
assert dep.get("owner") == ["repo", "organization"]
assert dep.get("repo") == ["owner"]
# Normal property schemas are still present and valid.
assert params["properties"]["owner"] == {"type": "string"}
assert params["properties"]["repo"] == {"type": "string"}
assert params["properties"]["organization"] == {"type": "string"}
def test_dependent_required_does_not_mutate_original_input():
"""The original schema's dependentRequired must be unchanged after sanitize."""
original_dep = {"owner": ["repo", "organization"], "repo": ["owner"]}
schema = {
"type": "object",
"properties": {
"owner": {"type": "string"},
"repo": {"type": "string"},
"organization": {"type": "string"},
},
"dependentRequired": {k: list(v) for k, v in original_dep.items()},
}
saved_copy = copy.deepcopy(schema)
tools = [_tool("t", schema)]
_ = sanitize_tool_schemas(tools)
assert schema == saved_copy
assert schema["dependentRequired"] == original_dep
def test_dependent_schemas_still_recursively_sanitized():
"""dependentSchemas (real schemas, not literal lists) must still be sanitized."""
schema = {
"type": "object",
"properties": {
"owner": {"type": "string"},
},
"dependentSchemas": {
"owner": {"type": "object"}, # bare object -- needs properties: {}
},
}
tools = [_tool("t", copy.deepcopy(schema))]
out = sanitize_tool_schemas(tools)
dep_schemas = out[0]["function"]["parameters"]["dependentSchemas"]
assert dep_schemas["owner"] == {"type": "object", "properties": {}}, (
f"dependentSchemas['owner'] was not fully sanitized: {dep_schemas['owner']!r}"
)