feat(api): implement dynamic max tokens handling for various providers

- Added _max_tokens_param method in AIAgent to return appropriate max tokens parameter based on the provider (OpenAI vs. others). - Updated API calls in AIAgent to utilize the new max tokens handling. - Introduced auxiliary_max_tokens_param function in auxiliary_client for consistent max tokens management across auxiliary clients. - Refactored multiple tools to use auxiliary_max_tokens_param for improved compatibility with different models and providers.
2026-04-25 00:51:20 +00:00 · 2026-02-26 20:23:56 -08:00 · 2026-02-26 20:23:56 -08:00 · 58fce0a37b
commit 58fce0a37b
parent f0458ebdb8
7 changed files with 67 additions and 20 deletions
--- a/tools/web_tools.py
+++ b/tools/web_tools.py
@ -242,7 +242,7 @@ Create a markdown summary that captures all key information in a well-organized,
            if _aux_async_client is None:
                logger.warning("No auxiliary model available for web content processing")
                return None
-            from agent.auxiliary_client import get_auxiliary_extra_body
+            from agent.auxiliary_client import get_auxiliary_extra_body, auxiliary_max_tokens_param
            _extra = get_auxiliary_extra_body()
            response = await _aux_async_client.chat.completions.create(
                model=model,
@ -251,7 +251,7 @@ Create a markdown summary that captures all key information in a well-organized,
                    {"role": "user", "content": user_prompt}
                ],
                temperature=0.1,
-                max_tokens=max_tokens,
+                **auxiliary_max_tokens_param(max_tokens),
                **({} if not _extra else {"extra_body": _extra}),
            )
            return response.choices[0].message.content.strip()
@ -365,7 +365,7 @@ Create a single, unified markdown summary."""
                fallback = fallback[:max_output_size] + "\n\n[... truncated ...]"
            return fallback

-        from agent.auxiliary_client import get_auxiliary_extra_body
+        from agent.auxiliary_client import get_auxiliary_extra_body, auxiliary_max_tokens_param
        _extra = get_auxiliary_extra_body()
        response = await _aux_async_client.chat.completions.create(
            model=model,
@ -374,7 +374,7 @@ Create a single, unified markdown summary."""
                {"role": "user", "content": synthesis_prompt}
            ],
            temperature=0.1,
-            max_tokens=4000,
+            **auxiliary_max_tokens_param(4000),
            **({} if not _extra else {"extra_body": _extra}),
        )
        final_summary = response.choices[0].message.content.strip()