Manually reformatted files

2024-10-25 13:32:25 +05:30
parent 2401e21ef2
commit a157e8e0a2
11 changed files with 175 additions and 95 deletions
--- a/examples/graph_visual_with_html.py
+++ b/examples/graph_visual_with_html.py
@@ -3,7 +3,7 @@ from pyvis.network import Network
 import random

 # Load the GraphML file
-G = nx.read_graphml('./dickens/graph_chunk_entity_relation.graphml')
+G = nx.read_graphml("./dickens/graph_chunk_entity_relation.graphml")

 # Create a Pyvis network
 net = Network(notebook=True)
@@ -13,7 +13,7 @@ net.from_nx(G)

 # Add colors to nodes
 for node in net.nodes:
-    node['color'] = "#{:06x}".format(random.randint(0, 0xFFFFFF))
+    node["color"] = "#{:06x}".format(random.randint(0, 0xFFFFFF))

 # Save and display the network
-net.show('knowledge_graph.html')
+net.show("knowledge_graph.html")
--- a/examples/graph_visual_with_neo4j.py
+++ b/examples/graph_visual_with_neo4j.py
@@ -13,6 +13,7 @@ NEO4J_URI = "bolt://localhost:7687"
 NEO4J_USERNAME = "neo4j"
 NEO4J_PASSWORD = "your_password"

+
 def convert_xml_to_json(xml_path, output_path):
    """Converts XML file to JSON and saves the output."""
    if not os.path.exists(xml_path):
@@ -21,7 +22,7 @@ def convert_xml_to_json(xml_path, output_path):

    json_data = xml_to_json(xml_path)
    if json_data:
-        with open(output_path, 'w', encoding='utf-8') as f:
+        with open(output_path, "w", encoding="utf-8") as f:
            json.dump(json_data, f, ensure_ascii=False, indent=2)
        print(f"JSON file created: {output_path}")
        return json_data
@@ -29,16 +30,18 @@ def convert_xml_to_json(xml_path, output_path):
        print("Failed to create JSON data")
        return None

+
 def process_in_batches(tx, query, data, batch_size):
    """Process data in batches and execute the given query."""
    for i in range(0, len(data), batch_size):
-        batch = data[i:i + batch_size]
+        batch = data[i : i + batch_size]
        tx.run(query, {"nodes": batch} if "nodes" in query else {"edges": batch})

+
 def main():
    # Paths
-    xml_file = os.path.join(WORKING_DIR, 'graph_chunk_entity_relation.graphml')
-    json_file = os.path.join(WORKING_DIR, 'graph_data.json')
+    xml_file = os.path.join(WORKING_DIR, "graph_chunk_entity_relation.graphml")
+    json_file = os.path.join(WORKING_DIR, "graph_data.json")

    # Convert XML to JSON
    json_data = convert_xml_to_json(xml_file, json_file)
@@ -46,8 +49,8 @@ def main():
        return

    # Load nodes and edges
-    nodes = json_data.get('nodes', [])
-    edges = json_data.get('edges', [])
+    nodes = json_data.get("nodes", [])
+    edges = json_data.get("edges", [])

    # Neo4j queries
    create_nodes_query = """
@@ -56,8 +59,8 @@ def main():
    SET e.entity_type = node.entity_type,
        e.description = node.description,
        e.source_id = node.source_id,
-        e.displayName = node.id  
-    REMOVE e:Entity  
+        e.displayName = node.id
+    REMOVE e:Entity
    WITH e, node
    CALL apoc.create.addLabels(e, [node.entity_type]) YIELD node AS labeledNode
    RETURN count(*)
@@ -100,19 +103,24 @@ def main():
        # Execute queries in batches
        with driver.session() as session:
            # Insert nodes in batches
-            session.execute_write(process_in_batches, create_nodes_query, nodes, BATCH_SIZE_NODES)
+            session.execute_write(
+                process_in_batches, create_nodes_query, nodes, BATCH_SIZE_NODES
+            )

            # Insert edges in batches
-            session.execute_write(process_in_batches, create_edges_query, edges, BATCH_SIZE_EDGES)
+            session.execute_write(
+                process_in_batches, create_edges_query, edges, BATCH_SIZE_EDGES
+            )

            # Set displayName and labels
            session.run(set_displayname_and_labels_query)

    except Exception as e:
        print(f"Error occurred: {e}")
-    
+
    finally:
        driver.close()

+
 if __name__ == "__main__":
    main()
--- a/examples/lightrag_openai_compatible_demo.py
+++ b/examples/lightrag_openai_compatible_demo.py
@@ -52,6 +52,7 @@ async def test_funcs():

 # asyncio.run(test_funcs())

+
 async def main():
    try:
        embedding_dimension = await get_embedding_dim()
@@ -61,35 +62,47 @@ async def main():
            working_dir=WORKING_DIR,
            llm_model_func=llm_model_func,
            embedding_func=EmbeddingFunc(
-                embedding_dim=embedding_dimension, max_token_size=8192, func=embedding_func
+                embedding_dim=embedding_dimension,
+                max_token_size=8192,
+                func=embedding_func,
            ),
        )

-
        with open("./book.txt", "r", encoding="utf-8") as f:
            rag.insert(f.read())

        # Perform naive search
        print(
-            rag.query("What are the top themes in this story?", param=QueryParam(mode="naive"))
+            rag.query(
+                "What are the top themes in this story?", param=QueryParam(mode="naive")
+            )
        )

        # Perform local search
        print(
-            rag.query("What are the top themes in this story?", param=QueryParam(mode="local"))
+            rag.query(
+                "What are the top themes in this story?", param=QueryParam(mode="local")
+            )
        )

        # Perform global search
        print(
-            rag.query("What are the top themes in this story?", param=QueryParam(mode="global"))
+            rag.query(
+                "What are the top themes in this story?",
+                param=QueryParam(mode="global"),
+            )
        )

        # Perform hybrid search
        print(
-            rag.query("What are the top themes in this story?", param=QueryParam(mode="hybrid"))
+            rag.query(
+                "What are the top themes in this story?",
+                param=QueryParam(mode="hybrid"),
+            )
        )
    except Exception as e:
        print(f"An error occurred: {e}")

+
 if __name__ == "__main__":
-    asyncio.run(main())
+    asyncio.run(main())
--- a/examples/lightrag_siliconcloud_demo.py
+++ b/examples/lightrag_siliconcloud_demo.py
@@ -30,7 +30,7 @@ async def embedding_func(texts: list[str]) -> np.ndarray:
        texts,
        model="netease-youdao/bce-embedding-base_v1",
        api_key=os.getenv("SILICONFLOW_API_KEY"),
-        max_token_size=512
+        max_token_size=512,
    )


--- a/examples/vram_management_demo.py
+++ b/examples/vram_management_demo.py
@@ -27,11 +27,12 @@ rag = LightRAG(
 # Read all .txt files from the TEXT_FILES_DIR directory
 texts = []
 for filename in os.listdir(TEXT_FILES_DIR):
-    if filename.endswith('.txt'):
+    if filename.endswith(".txt"):
        file_path = os.path.join(TEXT_FILES_DIR, filename)
-        with open(file_path, 'r', encoding='utf-8') as file:
+        with open(file_path, "r", encoding="utf-8") as file:
            texts.append(file.read())

+
 # Batch insert texts into LightRAG with a retry mechanism
 def insert_texts_with_retry(rag, texts, retries=3, delay=5):
    for _ in range(retries):
@@ -39,37 +40,58 @@ def insert_texts_with_retry(rag, texts, retries=3, delay=5):
            rag.insert(texts)
            return
        except Exception as e:
-            print(f"Error occurred during insertion: {e}. Retrying in {delay} seconds...")
+            print(
+                f"Error occurred during insertion: {e}. Retrying in {delay} seconds..."
+            )
            time.sleep(delay)
    raise RuntimeError("Failed to insert texts after multiple retries.")

+
 insert_texts_with_retry(rag, texts)

 # Perform different types of queries and handle potential errors
 try:
-    print(rag.query("What are the top themes in this story?", param=QueryParam(mode="naive")))
+    print(
+        rag.query(
+            "What are the top themes in this story?", param=QueryParam(mode="naive")
+        )
+    )
 except Exception as e:
    print(f"Error performing naive search: {e}")

 try:
-    print(rag.query("What are the top themes in this story?", param=QueryParam(mode="local")))
+    print(
+        rag.query(
+            "What are the top themes in this story?", param=QueryParam(mode="local")
+        )
+    )
 except Exception as e:
    print(f"Error performing local search: {e}")

 try:
-    print(rag.query("What are the top themes in this story?", param=QueryParam(mode="global")))
+    print(
+        rag.query(
+            "What are the top themes in this story?", param=QueryParam(mode="global")
+        )
+    )
 except Exception as e:
    print(f"Error performing global search: {e}")

 try:
-    print(rag.query("What are the top themes in this story?", param=QueryParam(mode="hybrid")))
+    print(
+        rag.query(
+            "What are the top themes in this story?", param=QueryParam(mode="hybrid")
+        )
+    )
 except Exception as e:
    print(f"Error performing hybrid search: {e}")

+
 # Function to clear VRAM resources
 def clear_vram():
    os.system("sudo nvidia-smi --gpu-reset")

+
 # Regularly clear VRAM to prevent overflow
 clear_vram_interval = 3600  # Clear once every hour
 start_time = time.time()