feat: enable custom LLM support for Crew.test()

- Add llm parameter to Crew.test() that accepts string or LLM instance - Maintain backward compatibility with openai_model_name parameter - Update CrewEvaluator to handle any LLM implementation - Add comprehensive test coverage Fixes #2076 Co-Authored-By: Joe Moura <joao@crewai.com>
2026-05-03 16:22:49 +00:00 · 2025-02-09 22:12:04 +00:00
parent 409892d65f
commit f3a681c7d9
3 changed files with 91 additions and 13 deletions
--- a/src/crewai/crew.py
+++ b/src/crewai/crew.py
@@ -1075,19 +1075,31 @@ class Crew(BaseModel):
    def test(
        self,
        n_iterations: int,
+        llm: Optional[Union[str, LLM]] = None,
        openai_model_name: Optional[str] = None,
        inputs: Optional[Dict[str, Any]] = None,
    ) -> None:
-        """Test and evaluate the Crew with the given inputs for n iterations concurrently using concurrent.futures."""
+        """Test and evaluate the Crew with the given inputs for n iterations concurrently using concurrent.futures.
+        
+        Args:
+            n_iterations: Number of iterations to run
+            llm: LLM instance or model name to use for evaluation
+            openai_model_name: (Deprecated) OpenAI model name for backward compatibility
+            inputs: Optional inputs for the crew
+        """
        test_crew = self.copy()

+        # Handle backward compatibility
+        if openai_model_name:
+            llm = openai_model_name
+
        self._test_execution_span = test_crew._telemetry.test_execution_span(
            test_crew,
            n_iterations,
            inputs,
-            openai_model_name,  # type: ignore[arg-type]
-        )  # type: ignore[arg-type]
-        evaluator = CrewEvaluator(test_crew, openai_model_name)  # type: ignore[arg-type]
+            str(llm) if isinstance(llm, str) else (llm.model if llm else None),
+        )
+        evaluator = CrewEvaluator(test_crew, llm)

        for i in range(1, n_iterations + 1):
            evaluator.set_iteration(i)