trustgraph/tests/unit/test_retrieval/test_document_rag.py

"""
Tests for DocumentRAG retrieval implementation
"""

import pytest
from unittest.mock import MagicMock, AsyncMock

from trustgraph.retrieval.document_rag.document_rag import DocumentRag, Query


class TestDocumentRag:
    """Test cases for DocumentRag class"""

    def test_document_rag_initialization_with_defaults(self):
        """Test DocumentRag initialization with default verbose setting"""
        # Create mock clients
        mock_prompt_client = MagicMock()
        mock_embeddings_client = MagicMock()
        mock_doc_embeddings_client = MagicMock()
        
        # Initialize DocumentRag
        document_rag = DocumentRag(
            prompt_client=mock_prompt_client,
            embeddings_client=mock_embeddings_client,
            doc_embeddings_client=mock_doc_embeddings_client
        )
        
        # Verify initialization
        assert document_rag.prompt_client == mock_prompt_client
        assert document_rag.embeddings_client == mock_embeddings_client
        assert document_rag.doc_embeddings_client == mock_doc_embeddings_client
        assert document_rag.verbose is False  # Default value

    def test_document_rag_initialization_with_verbose(self):
        """Test DocumentRag initialization with verbose enabled"""
        # Create mock clients
        mock_prompt_client = MagicMock()
        mock_embeddings_client = MagicMock()
        mock_doc_embeddings_client = MagicMock()
        
        # Initialize DocumentRag with verbose=True
        document_rag = DocumentRag(
            prompt_client=mock_prompt_client,
            embeddings_client=mock_embeddings_client,
            doc_embeddings_client=mock_doc_embeddings_client,
            verbose=True
        )
        
        # Verify initialization
        assert document_rag.prompt_client == mock_prompt_client
        assert document_rag.embeddings_client == mock_embeddings_client
        assert document_rag.doc_embeddings_client == mock_doc_embeddings_client
        assert document_rag.verbose is True


class TestQuery:
    """Test cases for Query class"""

    def test_query_initialization_with_defaults(self):
        """Test Query initialization with default parameters"""
        # Create mock DocumentRag
        mock_rag = MagicMock()
        
        # Initialize Query with defaults
        query = Query(
            rag=mock_rag,
            user="test_user",
            collection="test_collection",
            verbose=False
        )
        
        # Verify initialization
        assert query.rag == mock_rag
        assert query.user == "test_user"
        assert query.collection == "test_collection"
        assert query.verbose is False
        assert query.doc_limit == 20  # Default value

    def test_query_initialization_with_custom_doc_limit(self):
        """Test Query initialization with custom doc_limit"""
        # Create mock DocumentRag
        mock_rag = MagicMock()
        
        # Initialize Query with custom doc_limit
        query = Query(
            rag=mock_rag,
            user="custom_user",
            collection="custom_collection",
            verbose=True,
            doc_limit=50
        )
        
        # Verify initialization
        assert query.rag == mock_rag
        assert query.user == "custom_user"
        assert query.collection == "custom_collection"
        assert query.verbose is True
        assert query.doc_limit == 50

    @pytest.mark.asyncio
    async def test_get_vector_method(self):
        """Test Query.get_vector method calls embeddings client correctly"""
        # Create mock DocumentRag with embeddings client
        mock_rag = MagicMock()
        mock_embeddings_client = AsyncMock()
        mock_rag.embeddings_client = mock_embeddings_client
        
        # Mock the embed method to return test vectors
        expected_vectors = [[0.1, 0.2, 0.3], [0.4, 0.5, 0.6]]
        mock_embeddings_client.embed.return_value = expected_vectors
        
        # Initialize Query
        query = Query(
            rag=mock_rag,
            user="test_user",
            collection="test_collection",
            verbose=False
        )
        
        # Call get_vector
        test_query = "What documents are relevant?"
        result = await query.get_vector(test_query)
        
        # Verify embeddings client was called correctly
        mock_embeddings_client.embed.assert_called_once_with(test_query)
        
        # Verify result matches expected vectors
        assert result == expected_vectors

    @pytest.mark.asyncio
    async def test_get_docs_method(self):
        """Test Query.get_docs method retrieves documents correctly"""
        # Create mock DocumentRag with clients
        mock_rag = MagicMock()
        mock_embeddings_client = AsyncMock()
        mock_doc_embeddings_client = AsyncMock()
        mock_rag.embeddings_client = mock_embeddings_client
        mock_rag.doc_embeddings_client = mock_doc_embeddings_client
        
        # Mock the embedding and document query responses
        test_vectors = [[0.1, 0.2, 0.3]]
        mock_embeddings_client.embed.return_value = test_vectors
        
        # Mock document results
        test_docs = ["Document 1 content", "Document 2 content"]
        mock_doc_embeddings_client.query.return_value = test_docs
        
        # Initialize Query
        query = Query(
            rag=mock_rag,
            user="test_user",
            collection="test_collection",
            verbose=False,
            doc_limit=15
        )
        
        # Call get_docs
        test_query = "Find relevant documents"
        result = await query.get_docs(test_query)
        
        # Verify embeddings client was called
        mock_embeddings_client.embed.assert_called_once_with(test_query)
        
        # Verify doc embeddings client was called correctly
        mock_doc_embeddings_client.query.assert_called_once_with(
            test_vectors,
            limit=15,
            user="test_user",
            collection="test_collection"
        )
        
        # Verify result is list of documents
        assert result == test_docs

    @pytest.mark.asyncio
    async def test_document_rag_query_method(self):
        """Test DocumentRag.query method orchestrates full document RAG pipeline"""
        # Create mock clients
        mock_prompt_client = AsyncMock()
        mock_embeddings_client = AsyncMock()
        mock_doc_embeddings_client = AsyncMock()
        
        # Mock embeddings and document responses
        test_vectors = [[0.1, 0.2, 0.3]]
        test_docs = ["Relevant document content", "Another document"]
        expected_response = "This is the document RAG response"
        
        mock_embeddings_client.embed.return_value = test_vectors
        mock_doc_embeddings_client.query.return_value = test_docs
        mock_prompt_client.document_prompt.return_value = expected_response
        
        # Initialize DocumentRag
        document_rag = DocumentRag(
            prompt_client=mock_prompt_client,
            embeddings_client=mock_embeddings_client,
            doc_embeddings_client=mock_doc_embeddings_client,
            verbose=False
        )
        
        # Call DocumentRag.query
        result = await document_rag.query(
            query="test query",
            user="test_user",
            collection="test_collection",
            doc_limit=10
        )
        
        # Verify embeddings client was called
        mock_embeddings_client.embed.assert_called_once_with("test query")
        
        # Verify doc embeddings client was called
        mock_doc_embeddings_client.query.assert_called_once_with(
            test_vectors,
            limit=10,
            user="test_user",
            collection="test_collection"
        )
        
        # Verify prompt client was called with documents and query
        mock_prompt_client.document_prompt.assert_called_once_with(
            query="test query",
            documents=test_docs
        )
        
        # Verify result
        assert result == expected_response

    @pytest.mark.asyncio
    async def test_document_rag_query_with_defaults(self):
        """Test DocumentRag.query method with default parameters"""
        # Create mock clients
        mock_prompt_client = AsyncMock()
        mock_embeddings_client = AsyncMock()
        mock_doc_embeddings_client = AsyncMock()
        
        # Mock responses
        mock_embeddings_client.embed.return_value = [[0.1, 0.2]]
        mock_doc_embeddings_client.query.return_value = ["Default doc"]
        mock_prompt_client.document_prompt.return_value = "Default response"
        
        # Initialize DocumentRag
        document_rag = DocumentRag(
            prompt_client=mock_prompt_client,
            embeddings_client=mock_embeddings_client,
            doc_embeddings_client=mock_doc_embeddings_client
        )
        
        # Call DocumentRag.query with minimal parameters
        result = await document_rag.query("simple query")
        
        # Verify default parameters were used
        mock_doc_embeddings_client.query.assert_called_once_with(
            [[0.1, 0.2]],
            limit=20,  # Default doc_limit
            user="trustgraph",  # Default user
            collection="default"  # Default collection
        )
        
        assert result == "Default response"

    @pytest.mark.asyncio
    async def test_get_docs_with_verbose_output(self):
        """Test Query.get_docs method with verbose logging"""
        # Create mock DocumentRag with clients
        mock_rag = MagicMock()
        mock_embeddings_client = AsyncMock()
        mock_doc_embeddings_client = AsyncMock()
        mock_rag.embeddings_client = mock_embeddings_client
        mock_rag.doc_embeddings_client = mock_doc_embeddings_client
        
        # Mock responses
        mock_embeddings_client.embed.return_value = [[0.7, 0.8]]
        mock_doc_embeddings_client.query.return_value = ["Verbose test doc"]
        
        # Initialize Query with verbose=True
        query = Query(
            rag=mock_rag,
            user="test_user",
            collection="test_collection",
            verbose=True,
            doc_limit=5
        )
        
        # Call get_docs
        result = await query.get_docs("verbose test")
        
        # Verify calls were made
        mock_embeddings_client.embed.assert_called_once_with("verbose test")
        mock_doc_embeddings_client.query.assert_called_once()
        
        # Verify result
        assert result == ["Verbose test doc"]

    @pytest.mark.asyncio
    async def test_document_rag_query_with_verbose(self):
        """Test DocumentRag.query method with verbose logging enabled"""
        # Create mock clients
        mock_prompt_client = AsyncMock()
        mock_embeddings_client = AsyncMock()
        mock_doc_embeddings_client = AsyncMock()
        
        # Mock responses
        mock_embeddings_client.embed.return_value = [[0.3, 0.4]]
        mock_doc_embeddings_client.query.return_value = ["Verbose doc content"]
        mock_prompt_client.document_prompt.return_value = "Verbose RAG response"
        
        # Initialize DocumentRag with verbose=True
        document_rag = DocumentRag(
            prompt_client=mock_prompt_client,
            embeddings_client=mock_embeddings_client,
            doc_embeddings_client=mock_doc_embeddings_client,
            verbose=True
        )
        
        # Call DocumentRag.query
        result = await document_rag.query("verbose query test")
        
        # Verify all clients were called
        mock_embeddings_client.embed.assert_called_once_with("verbose query test")
        mock_doc_embeddings_client.query.assert_called_once()
        mock_prompt_client.document_prompt.assert_called_once_with(
            query="verbose query test",
            documents=["Verbose doc content"]
        )
        
        assert result == "Verbose RAG response"

    @pytest.mark.asyncio
    async def test_get_docs_with_empty_results(self):
        """Test Query.get_docs method when no documents are found"""
        # Create mock DocumentRag with clients
        mock_rag = MagicMock()
        mock_embeddings_client = AsyncMock()
        mock_doc_embeddings_client = AsyncMock()
        mock_rag.embeddings_client = mock_embeddings_client
        mock_rag.doc_embeddings_client = mock_doc_embeddings_client
        
        # Mock responses - empty document list
        mock_embeddings_client.embed.return_value = [[0.1, 0.2]]
        mock_doc_embeddings_client.query.return_value = []  # No documents found
        
        # Initialize Query
        query = Query(
            rag=mock_rag,
            user="test_user",
            collection="test_collection",
            verbose=False
        )
        
        # Call get_docs
        result = await query.get_docs("query with no results")
        
        # Verify calls were made
        mock_embeddings_client.embed.assert_called_once_with("query with no results")
        mock_doc_embeddings_client.query.assert_called_once()
        
        # Verify empty result is returned
        assert result == []

    @pytest.mark.asyncio
    async def test_document_rag_query_with_empty_documents(self):
        """Test DocumentRag.query method when no documents are retrieved"""
        # Create mock clients
        mock_prompt_client = AsyncMock()
        mock_embeddings_client = AsyncMock()
        mock_doc_embeddings_client = AsyncMock()
        
        # Mock responses - no documents found
        mock_embeddings_client.embed.return_value = [[0.5, 0.6]]
        mock_doc_embeddings_client.query.return_value = []  # Empty document list
        mock_prompt_client.document_prompt.return_value = "No documents found response"
        
        # Initialize DocumentRag
        document_rag = DocumentRag(
            prompt_client=mock_prompt_client,
            embeddings_client=mock_embeddings_client,
            doc_embeddings_client=mock_doc_embeddings_client,
            verbose=False
        )
        
        # Call DocumentRag.query
        result = await document_rag.query("query with no matching docs")
        
        # Verify prompt client was called with empty document list
        mock_prompt_client.document_prompt.assert_called_once_with(
            query="query with no matching docs",
            documents=[]
        )
        
        assert result == "No documents found response"

    @pytest.mark.asyncio
    async def test_get_vector_with_verbose(self):
        """Test Query.get_vector method with verbose logging"""
        # Create mock DocumentRag with embeddings client
        mock_rag = MagicMock()
        mock_embeddings_client = AsyncMock()
        mock_rag.embeddings_client = mock_embeddings_client
        
        # Mock the embed method
        expected_vectors = [[0.9, 1.0, 1.1]]
        mock_embeddings_client.embed.return_value = expected_vectors
        
        # Initialize Query with verbose=True
        query = Query(
            rag=mock_rag,
            user="test_user",
            collection="test_collection",
            verbose=True
        )
        
        # Call get_vector
        result = await query.get_vector("verbose vector test")
        
        # Verify embeddings client was called
        mock_embeddings_client.embed.assert_called_once_with("verbose vector test")
        
        # Verify result
        assert result == expected_vectors

    @pytest.mark.asyncio
    async def test_document_rag_integration_flow(self):
        """Test complete DocumentRag integration with realistic data flow"""
        # Create mock clients
        mock_prompt_client = AsyncMock()
        mock_embeddings_client = AsyncMock()
        mock_doc_embeddings_client = AsyncMock()
        
        # Mock realistic responses
        query_text = "What is machine learning?"
        query_vectors = [[0.1, 0.2, 0.3, 0.4, 0.5]]
        retrieved_docs = [
            "Machine learning is a subset of artificial intelligence...",
            "ML algorithms learn patterns from data to make predictions...",
            "Common ML techniques include supervised and unsupervised learning..."
        ]
        final_response = "Machine learning is a field of AI that enables computers to learn and improve from experience without being explicitly programmed."
        
        mock_embeddings_client.embed.return_value = query_vectors
        mock_doc_embeddings_client.query.return_value = retrieved_docs
        mock_prompt_client.document_prompt.return_value = final_response
        
        # Initialize DocumentRag
        document_rag = DocumentRag(
            prompt_client=mock_prompt_client,
            embeddings_client=mock_embeddings_client,
            doc_embeddings_client=mock_doc_embeddings_client,
            verbose=False
        )
        
        # Execute full pipeline
        result = await document_rag.query(
            query=query_text,
            user="research_user", 
            collection="ml_knowledge",
            doc_limit=25
        )
        
        # Verify complete pipeline execution
        mock_embeddings_client.embed.assert_called_once_with(query_text)
        
        mock_doc_embeddings_client.query.assert_called_once_with(
            query_vectors,
            limit=25,
            user="research_user",
            collection="ml_knowledge"
        )
        
        mock_prompt_client.document_prompt.assert_called_once_with(
            query=query_text,
            documents=retrieved_docs
        )
        
        # Verify final result
        assert result == final_response
Release/v1.2 (#457) * Bump setup.py versions for 1.1 * PoC MCP server (#419) * Very initial MCP server PoC for TrustGraph * Put service on port 8000 * Add MCP container and packages to buildout * Update docs for API/CLI changes in 1.0 (#421) * Update some API basics for the 0.23/1.0 API change * Add MCP container push (#425) * Add command args to the MCP server (#426) * Host and port parameters * Added websocket arg * More docs * MCP client support (#427) - MCP client service - Tool request/response schema - API gateway support for mcp-tool - Message translation for tool request & response - Make mcp-tool using configuration service for information about where the MCP services are. * Feature/react call mcp (#428) Key Features - MCP Tool Integration: Added core MCP tool support with ToolClientSpec and ToolClient classes - API Enhancement: New mcp_tool method for flow-specific tool invocation - CLI Tooling: New tg-invoke-mcp-tool command for testing MCP integration - React Agent Enhancement: Fixed and improved multi-tool invocation capabilities - Tool Management: Enhanced CLI for tool configuration and management Changes - Added MCP tool invocation to API with flow-specific integration - Implemented ToolClientSpec and ToolClient for tool call handling - Updated agent-manager-react to invoke MCP tools with configurable types - Enhanced CLI with new commands and improved help text - Added comprehensive documentation for new CLI commands - Improved tool configuration management Testing - Added tg-invoke-mcp-tool CLI command for isolated MCP integration testing - Enhanced agent capability to invoke multiple tools simultaneously * Test suite executed from CI pipeline (#433) * Test strategy & test cases * Unit tests * Integration tests * Extending test coverage (#434) * Contract tests * Testing embeedings * Agent unit tests * Knowledge pipeline tests * Turn on contract tests * Increase storage test coverage (#435) * Fixing storage and adding tests * PR pipeline only runs quick tests * Empty configuration is returned as empty list, previously was not in response (#436) * Update config util to take files as well as command-line text (#437) * Updated CLI invocation and config model for tools and mcp (#438) * Updated CLI invocation and config model for tools and mcp * CLI anomalies * Tweaked the MCP tool implementation for new model * Update agent implementation to match the new model * Fix agent tools, now all tested * Fixed integration tests * Fix MCP delete tool params * Update Python deps to 1.2 * Update to enable knowledge extraction using the agent framework (#439) * Implement KG extraction agent (kg-extract-agent) * Using ReAct framework (agent-manager-react) * ReAct manager had an issue when emitting JSON, which conflicts which ReAct manager's own JSON messages, so refactored ReAct manager to use traditional ReAct messages, non-JSON structure. * Minor refactor to take the prompt template client out of prompt-template so it can be more readily used by other modules. kg-extract-agent uses this framework. * Migrate from setup.py to pyproject.toml (#440) * Converted setup.py to pyproject.toml * Modern package infrastructure as recommended by py docs * Install missing build deps (#441) * Install missing build deps (#442) * Implement logging strategy (#444) * Logging strategy and convert all prints() to logging invocations * Fix/startup failure (#445) * Fix loggin startup problems * Fix logging startup problems (#446) * Fix logging startup problems (#447) * Fixed Mistral OCR to use current API (#448) * Fixed Mistral OCR to use current API * Added PDF decoder tests * Fix Mistral OCR ident to be standard pdf-decoder (#450) * Fix Mistral OCR ident to be standard pdf-decoder * Correct test * Schema structure refactor (#451) * Write schema refactor spec * Implemented schema refactor spec * Structure data mvp (#452) * Structured data tech spec * Architecture principles * New schemas * Updated schemas and specs * Object extractor * Add .coveragerc * New tests * Cassandra object storage * Trying to object extraction working, issues exist * Validate librarian collection (#453) * Fix token chunker, broken API invocation (#454) * Fix token chunker, broken API invocation (#455) * Knowledge load utility CLI (#456) * Knowledge loader * More tests 2025-08-18 20:56:09 +01:00			`"""`
			`Tests for DocumentRAG retrieval implementation`
			`"""`

			`import pytest`
			`from unittest.mock import MagicMock, AsyncMock`

			`from trustgraph.retrieval.document_rag.document_rag import DocumentRag, Query`


			`class TestDocumentRag:`
			`"""Test cases for DocumentRag class"""`

			`def test_document_rag_initialization_with_defaults(self):`
			`"""Test DocumentRag initialization with default verbose setting"""`
			`# Create mock clients`
			`mock_prompt_client = MagicMock()`
			`mock_embeddings_client = MagicMock()`
			`mock_doc_embeddings_client = MagicMock()`

			`# Initialize DocumentRag`
			`document_rag = DocumentRag(`
			`prompt_client=mock_prompt_client,`
			`embeddings_client=mock_embeddings_client,`
			`doc_embeddings_client=mock_doc_embeddings_client`
			`)`

			`# Verify initialization`
			`assert document_rag.prompt_client == mock_prompt_client`
			`assert document_rag.embeddings_client == mock_embeddings_client`
			`assert document_rag.doc_embeddings_client == mock_doc_embeddings_client`
			`assert document_rag.verbose is False # Default value`

			`def test_document_rag_initialization_with_verbose(self):`
			`"""Test DocumentRag initialization with verbose enabled"""`
			`# Create mock clients`
			`mock_prompt_client = MagicMock()`
			`mock_embeddings_client = MagicMock()`
			`mock_doc_embeddings_client = MagicMock()`

			`# Initialize DocumentRag with verbose=True`
			`document_rag = DocumentRag(`
			`prompt_client=mock_prompt_client,`
			`embeddings_client=mock_embeddings_client,`
			`doc_embeddings_client=mock_doc_embeddings_client,`
			`verbose=True`
			`)`

			`# Verify initialization`
			`assert document_rag.prompt_client == mock_prompt_client`
			`assert document_rag.embeddings_client == mock_embeddings_client`
			`assert document_rag.doc_embeddings_client == mock_doc_embeddings_client`
			`assert document_rag.verbose is True`


			`class TestQuery:`
			`"""Test cases for Query class"""`

			`def test_query_initialization_with_defaults(self):`
			`"""Test Query initialization with default parameters"""`
			`# Create mock DocumentRag`
			`mock_rag = MagicMock()`

			`# Initialize Query with defaults`
			`query = Query(`
			`rag=mock_rag,`
			`user="test_user",`
			`collection="test_collection",`
			`verbose=False`
			`)`

			`# Verify initialization`
			`assert query.rag == mock_rag`
			`assert query.user == "test_user"`
			`assert query.collection == "test_collection"`
			`assert query.verbose is False`
			`assert query.doc_limit == 20 # Default value`

			`def test_query_initialization_with_custom_doc_limit(self):`
			`"""Test Query initialization with custom doc_limit"""`
			`# Create mock DocumentRag`
			`mock_rag = MagicMock()`

			`# Initialize Query with custom doc_limit`
			`query = Query(`
			`rag=mock_rag,`
			`user="custom_user",`
			`collection="custom_collection",`
			`verbose=True,`
			`doc_limit=50`
			`)`

			`# Verify initialization`
			`assert query.rag == mock_rag`
			`assert query.user == "custom_user"`
			`assert query.collection == "custom_collection"`
			`assert query.verbose is True`
			`assert query.doc_limit == 50`

			`@pytest.mark.asyncio`
			`async def test_get_vector_method(self):`
			`"""Test Query.get_vector method calls embeddings client correctly"""`
			`# Create mock DocumentRag with embeddings client`
			`mock_rag = MagicMock()`
			`mock_embeddings_client = AsyncMock()`
			`mock_rag.embeddings_client = mock_embeddings_client`

			`# Mock the embed method to return test vectors`
			`expected_vectors = [[0.1, 0.2, 0.3], [0.4, 0.5, 0.6]]`
			`mock_embeddings_client.embed.return_value = expected_vectors`

			`# Initialize Query`
			`query = Query(`
			`rag=mock_rag,`
			`user="test_user",`
			`collection="test_collection",`
			`verbose=False`
			`)`

			`# Call get_vector`
			`test_query = "What documents are relevant?"`
			`result = await query.get_vector(test_query)`

			`# Verify embeddings client was called correctly`
			`mock_embeddings_client.embed.assert_called_once_with(test_query)`

			`# Verify result matches expected vectors`
			`assert result == expected_vectors`

			`@pytest.mark.asyncio`
			`async def test_get_docs_method(self):`
			`"""Test Query.get_docs method retrieves documents correctly"""`
			`# Create mock DocumentRag with clients`
			`mock_rag = MagicMock()`
			`mock_embeddings_client = AsyncMock()`
			`mock_doc_embeddings_client = AsyncMock()`
			`mock_rag.embeddings_client = mock_embeddings_client`
			`mock_rag.doc_embeddings_client = mock_doc_embeddings_client`

			`# Mock the embedding and document query responses`
			`test_vectors = [[0.1, 0.2, 0.3]]`
			`mock_embeddings_client.embed.return_value = test_vectors`

			`# Mock document results`
			`test_docs = ["Document 1 content", "Document 2 content"]`
			`mock_doc_embeddings_client.query.return_value = test_docs`

			`# Initialize Query`
			`query = Query(`
			`rag=mock_rag,`
			`user="test_user",`
			`collection="test_collection",`
			`verbose=False,`
			`doc_limit=15`
			`)`

			`# Call get_docs`
			`test_query = "Find relevant documents"`
			`result = await query.get_docs(test_query)`

			`# Verify embeddings client was called`
			`mock_embeddings_client.embed.assert_called_once_with(test_query)`

			`# Verify doc embeddings client was called correctly`
			`mock_doc_embeddings_client.query.assert_called_once_with(`
			`test_vectors,`
			`limit=15,`
			`user="test_user",`
			`collection="test_collection"`
			`)`

			`# Verify result is list of documents`
			`assert result == test_docs`

			`@pytest.mark.asyncio`
			`async def test_document_rag_query_method(self):`
			`"""Test DocumentRag.query method orchestrates full document RAG pipeline"""`
			`# Create mock clients`
			`mock_prompt_client = AsyncMock()`
			`mock_embeddings_client = AsyncMock()`
			`mock_doc_embeddings_client = AsyncMock()`

			`# Mock embeddings and document responses`
			`test_vectors = [[0.1, 0.2, 0.3]]`
			`test_docs = ["Relevant document content", "Another document"]`
			`expected_response = "This is the document RAG response"`

			`mock_embeddings_client.embed.return_value = test_vectors`
			`mock_doc_embeddings_client.query.return_value = test_docs`
			`mock_prompt_client.document_prompt.return_value = expected_response`

			`# Initialize DocumentRag`
			`document_rag = DocumentRag(`
			`prompt_client=mock_prompt_client,`
			`embeddings_client=mock_embeddings_client,`
			`doc_embeddings_client=mock_doc_embeddings_client,`
			`verbose=False`
			`)`

			`# Call DocumentRag.query`
			`result = await document_rag.query(`
			`query="test query",`
			`user="test_user",`
			`collection="test_collection",`
			`doc_limit=10`
			`)`

			`# Verify embeddings client was called`
			`mock_embeddings_client.embed.assert_called_once_with("test query")`

			`# Verify doc embeddings client was called`
			`mock_doc_embeddings_client.query.assert_called_once_with(`
			`test_vectors,`
			`limit=10,`
			`user="test_user",`
			`collection="test_collection"`
			`)`

			`# Verify prompt client was called with documents and query`
			`mock_prompt_client.document_prompt.assert_called_once_with(`
			`query="test query",`
			`documents=test_docs`
			`)`

			`# Verify result`
			`assert result == expected_response`

			`@pytest.mark.asyncio`
			`async def test_document_rag_query_with_defaults(self):`
			`"""Test DocumentRag.query method with default parameters"""`
			`# Create mock clients`
			`mock_prompt_client = AsyncMock()`
			`mock_embeddings_client = AsyncMock()`
			`mock_doc_embeddings_client = AsyncMock()`

			`# Mock responses`
			`mock_embeddings_client.embed.return_value = [[0.1, 0.2]]`
			`mock_doc_embeddings_client.query.return_value = ["Default doc"]`
			`mock_prompt_client.document_prompt.return_value = "Default response"`

			`# Initialize DocumentRag`
			`document_rag = DocumentRag(`
			`prompt_client=mock_prompt_client,`
			`embeddings_client=mock_embeddings_client,`
			`doc_embeddings_client=mock_doc_embeddings_client`
			`)`

			`# Call DocumentRag.query with minimal parameters`
			`result = await document_rag.query("simple query")`

			`# Verify default parameters were used`
			`mock_doc_embeddings_client.query.assert_called_once_with(`
			`[[0.1, 0.2]],`
			`limit=20, # Default doc_limit`
			`user="trustgraph", # Default user`
			`collection="default" # Default collection`
			`)`

			`assert result == "Default response"`

			`@pytest.mark.asyncio`
			`async def test_get_docs_with_verbose_output(self):`
			`"""Test Query.get_docs method with verbose logging"""`
			`# Create mock DocumentRag with clients`
			`mock_rag = MagicMock()`
			`mock_embeddings_client = AsyncMock()`
			`mock_doc_embeddings_client = AsyncMock()`
			`mock_rag.embeddings_client = mock_embeddings_client`
			`mock_rag.doc_embeddings_client = mock_doc_embeddings_client`

			`# Mock responses`
			`mock_embeddings_client.embed.return_value = [[0.7, 0.8]]`
			`mock_doc_embeddings_client.query.return_value = ["Verbose test doc"]`

			`# Initialize Query with verbose=True`
			`query = Query(`
			`rag=mock_rag,`
			`user="test_user",`
			`collection="test_collection",`
			`verbose=True,`
			`doc_limit=5`
			`)`

			`# Call get_docs`
			`result = await query.get_docs("verbose test")`

			`# Verify calls were made`
			`mock_embeddings_client.embed.assert_called_once_with("verbose test")`
			`mock_doc_embeddings_client.query.assert_called_once()`

			`# Verify result`
			`assert result == ["Verbose test doc"]`

			`@pytest.mark.asyncio`
			`async def test_document_rag_query_with_verbose(self):`
			`"""Test DocumentRag.query method with verbose logging enabled"""`
			`# Create mock clients`
			`mock_prompt_client = AsyncMock()`
			`mock_embeddings_client = AsyncMock()`
			`mock_doc_embeddings_client = AsyncMock()`

			`# Mock responses`
			`mock_embeddings_client.embed.return_value = [[0.3, 0.4]]`
			`mock_doc_embeddings_client.query.return_value = ["Verbose doc content"]`
			`mock_prompt_client.document_prompt.return_value = "Verbose RAG response"`

			`# Initialize DocumentRag with verbose=True`
			`document_rag = DocumentRag(`
			`prompt_client=mock_prompt_client,`
			`embeddings_client=mock_embeddings_client,`
			`doc_embeddings_client=mock_doc_embeddings_client,`
			`verbose=True`
			`)`

			`# Call DocumentRag.query`
			`result = await document_rag.query("verbose query test")`

			`# Verify all clients were called`
			`mock_embeddings_client.embed.assert_called_once_with("verbose query test")`
			`mock_doc_embeddings_client.query.assert_called_once()`
			`mock_prompt_client.document_prompt.assert_called_once_with(`
			`query="verbose query test",`
			`documents=["Verbose doc content"]`
			`)`

			`assert result == "Verbose RAG response"`

			`@pytest.mark.asyncio`
			`async def test_get_docs_with_empty_results(self):`
			`"""Test Query.get_docs method when no documents are found"""`
			`# Create mock DocumentRag with clients`
			`mock_rag = MagicMock()`
			`mock_embeddings_client = AsyncMock()`
			`mock_doc_embeddings_client = AsyncMock()`
			`mock_rag.embeddings_client = mock_embeddings_client`
			`mock_rag.doc_embeddings_client = mock_doc_embeddings_client`

			`# Mock responses - empty document list`
			`mock_embeddings_client.embed.return_value = [[0.1, 0.2]]`
			`mock_doc_embeddings_client.query.return_value = [] # No documents found`

			`# Initialize Query`
			`query = Query(`
			`rag=mock_rag,`
			`user="test_user",`
			`collection="test_collection",`
			`verbose=False`
			`)`

			`# Call get_docs`
			`result = await query.get_docs("query with no results")`

			`# Verify calls were made`
			`mock_embeddings_client.embed.assert_called_once_with("query with no results")`
			`mock_doc_embeddings_client.query.assert_called_once()`

			`# Verify empty result is returned`
			`assert result == []`

			`@pytest.mark.asyncio`
			`async def test_document_rag_query_with_empty_documents(self):`
			`"""Test DocumentRag.query method when no documents are retrieved"""`
			`# Create mock clients`
			`mock_prompt_client = AsyncMock()`
			`mock_embeddings_client = AsyncMock()`
			`mock_doc_embeddings_client = AsyncMock()`

			`# Mock responses - no documents found`
			`mock_embeddings_client.embed.return_value = [[0.5, 0.6]]`
			`mock_doc_embeddings_client.query.return_value = [] # Empty document list`
			`mock_prompt_client.document_prompt.return_value = "No documents found response"`

			`# Initialize DocumentRag`
			`document_rag = DocumentRag(`
			`prompt_client=mock_prompt_client,`
			`embeddings_client=mock_embeddings_client,`
			`doc_embeddings_client=mock_doc_embeddings_client,`
			`verbose=False`
			`)`

			`# Call DocumentRag.query`
			`result = await document_rag.query("query with no matching docs")`

			`# Verify prompt client was called with empty document list`
			`mock_prompt_client.document_prompt.assert_called_once_with(`
			`query="query with no matching docs",`
			`documents=[]`
			`)`

			`assert result == "No documents found response"`

			`@pytest.mark.asyncio`
			`async def test_get_vector_with_verbose(self):`
			`"""Test Query.get_vector method with verbose logging"""`
			`# Create mock DocumentRag with embeddings client`
			`mock_rag = MagicMock()`
			`mock_embeddings_client = AsyncMock()`
			`mock_rag.embeddings_client = mock_embeddings_client`

			`# Mock the embed method`
			`expected_vectors = [[0.9, 1.0, 1.1]]`
			`mock_embeddings_client.embed.return_value = expected_vectors`

			`# Initialize Query with verbose=True`
			`query = Query(`
			`rag=mock_rag,`
			`user="test_user",`
			`collection="test_collection",`
			`verbose=True`
			`)`

			`# Call get_vector`
			`result = await query.get_vector("verbose vector test")`

			`# Verify embeddings client was called`
			`mock_embeddings_client.embed.assert_called_once_with("verbose vector test")`

			`# Verify result`
			`assert result == expected_vectors`

			`@pytest.mark.asyncio`
			`async def test_document_rag_integration_flow(self):`
			`"""Test complete DocumentRag integration with realistic data flow"""`
			`# Create mock clients`
			`mock_prompt_client = AsyncMock()`
			`mock_embeddings_client = AsyncMock()`
			`mock_doc_embeddings_client = AsyncMock()`

			`# Mock realistic responses`
			`query_text = "What is machine learning?"`
			`query_vectors = [[0.1, 0.2, 0.3, 0.4, 0.5]]`
			`retrieved_docs = [`
			`"Machine learning is a subset of artificial intelligence...",`
			`"ML algorithms learn patterns from data to make predictions...",`
			`"Common ML techniques include supervised and unsupervised learning..."`
			`]`
			`final_response = "Machine learning is a field of AI that enables computers to learn and improve from experience without being explicitly programmed."`

			`mock_embeddings_client.embed.return_value = query_vectors`
			`mock_doc_embeddings_client.query.return_value = retrieved_docs`
			`mock_prompt_client.document_prompt.return_value = final_response`

			`# Initialize DocumentRag`
			`document_rag = DocumentRag(`
			`prompt_client=mock_prompt_client,`
			`embeddings_client=mock_embeddings_client,`
			`doc_embeddings_client=mock_doc_embeddings_client,`
			`verbose=False`
			`)`

			`# Execute full pipeline`
			`result = await document_rag.query(`
			`query=query_text,`
			`user="research_user",`
			`collection="ml_knowledge",`
			`doc_limit=25`
			`)`

			`# Verify complete pipeline execution`
			`mock_embeddings_client.embed.assert_called_once_with(query_text)`

			`mock_doc_embeddings_client.query.assert_called_once_with(`
			`query_vectors,`
			`limit=25,`
			`user="research_user",`
			`collection="ml_knowledge"`
			`)`

			`mock_prompt_client.document_prompt.assert_called_once_with(`
			`query=query_text,`
			`documents=retrieved_docs`
			`)`

			`# Verify final result`
			`assert result == final_response`