mirror of
https://github.com/trustgraph-ai/trustgraph.git
synced 2026-07-25 05:01:01 +02:00
Logging strategy updates
This commit is contained in:
parent
0ad66dffb2
commit
a520478af1
2 changed files with 10 additions and 6 deletions
|
|
@ -517,7 +517,7 @@ class McpServer:
|
||||||
|
|
||||||
async for response in gen:
|
async for response in gen:
|
||||||
|
|
||||||
print(response)
|
logging.debug(f"Agent response: {response}")
|
||||||
|
|
||||||
if "thought" in response:
|
if "thought" in response:
|
||||||
await ctx.session.send_log_message(
|
await ctx.session.send_log_message(
|
||||||
|
|
|
||||||
|
|
@ -6,12 +6,16 @@ PDF document as text as separate output objects.
|
||||||
|
|
||||||
import tempfile
|
import tempfile
|
||||||
import base64
|
import base64
|
||||||
|
import logging
|
||||||
import pytesseract
|
import pytesseract
|
||||||
from pdf2image import convert_from_bytes
|
from pdf2image import convert_from_bytes
|
||||||
|
|
||||||
from ... schema import Document, TextDocument, Metadata
|
from ... schema import Document, TextDocument, Metadata
|
||||||
from ... base import FlowProcessor, ConsumerSpec, ProducerSpec
|
from ... base import FlowProcessor, ConsumerSpec, ProducerSpec
|
||||||
|
|
||||||
|
# Module logger
|
||||||
|
logger = logging.getLogger(__name__)
|
||||||
|
|
||||||
default_ident = "pdf-decoder"
|
default_ident = "pdf-decoder"
|
||||||
|
|
||||||
class Processor(FlowProcessor):
|
class Processor(FlowProcessor):
|
||||||
|
|
@ -41,15 +45,15 @@ class Processor(FlowProcessor):
|
||||||
)
|
)
|
||||||
)
|
)
|
||||||
|
|
||||||
print("PDF OCR inited")
|
logger.info("PDF OCR processor initialized")
|
||||||
|
|
||||||
async def on_message(self, msg, consumer, flow):
|
async def on_message(self, msg, consumer, flow):
|
||||||
|
|
||||||
print("PDF message received", flush=True)
|
logger.info("PDF message received")
|
||||||
|
|
||||||
v = msg.value()
|
v = msg.value()
|
||||||
|
|
||||||
print(f"Decoding {v.metadata.id}...", flush=True)
|
logger.info(f"Decoding {v.metadata.id}...")
|
||||||
|
|
||||||
blob = base64.b64decode(v.data)
|
blob = base64.b64decode(v.data)
|
||||||
|
|
||||||
|
|
@ -60,7 +64,7 @@ class Processor(FlowProcessor):
|
||||||
try:
|
try:
|
||||||
text = pytesseract.image_to_string(page, lang='eng')
|
text = pytesseract.image_to_string(page, lang='eng')
|
||||||
except Exception as e:
|
except Exception as e:
|
||||||
print(f"Page did not OCR: {e}")
|
logger.warning(f"Page did not OCR: {e}")
|
||||||
continue
|
continue
|
||||||
|
|
||||||
r = TextDocument(
|
r = TextDocument(
|
||||||
|
|
@ -70,7 +74,7 @@ class Processor(FlowProcessor):
|
||||||
|
|
||||||
await flow("output").send(r)
|
await flow("output").send(r)
|
||||||
|
|
||||||
print("Done.", flush=True)
|
logger.info("PDF decoding complete")
|
||||||
|
|
||||||
@staticmethod
|
@staticmethod
|
||||||
def add_args(parser):
|
def add_args(parser):
|
||||||
|
|
|
||||||
Loading…
Add table
Add a link
Reference in a new issue