From a6765659cbeffeae48111f0797d3b4d0752ae80c Mon Sep 17 00:00:00 2001 From: Ellie Date: Tue, 27 Aug 2024 12:02:19 +0100 Subject: add test progress --- tests/test_load_lambda.py | 7 ++++--- 1 file changed, 4 insertions(+), 3 deletions(-) (limited to 'tests/test_load_lambda.py') diff --git a/tests/test_load_lambda.py b/tests/test_load_lambda.py index 88c71e4..30e55f3 100644 --- a/tests/test_load_lambda.py +++ b/tests/test_load_lambda.py @@ -18,7 +18,7 @@ from src.load_lambda import ( def aws_credentials(): os.environ["AWS_ACCESS_KEY_ID"] = "testing" os.environ["AWS_SECRET_ACCESS_KEY"] = "testing" - os.environ["AWS_SECURIT_TOKEN"] = "testing" + os.environ["AWS_SECURITY_TOKEN"] = "testing" os.environ["AWS_SESSION_TOKEN"] = "testing" os.environ["AWS_DEFAULT_REGION"] = "eu-west-2" @@ -88,5 +88,6 @@ class TestConvertParquetToDfs: # assert "dim_staff" in result -class TestUploadDfsToDatabase: - pass +@pytest.fixture +def mock_parquet_file(mocker): + return mocker.patch(src.load_lambda.convert_parquet_files_to_dfs()) -- cgit v1.2.3 From ec4a953ac73e6b828c61defe4d234a690461fcb6 Mon Sep 17 00:00:00 2001 From: Ellie Date: Tue, 27 Aug 2024 12:28:27 +0100 Subject: add first retrieve secrets test --- tests/test_load_lambda.py | 44 +++++++++++++++++++++++++++++++++----------- 1 file changed, 33 insertions(+), 11 deletions(-) (limited to 'tests/test_load_lambda.py') diff --git a/tests/test_load_lambda.py b/tests/test_load_lambda.py index 30e55f3..3df94e4 100644 --- a/tests/test_load_lambda.py +++ b/tests/test_load_lambda.py @@ -5,13 +5,7 @@ from moto import mock_aws import boto3 import os import pytest -from src.load_lambda import ( - lambda_handler, - connect_to_db_and_return_engine, - get_transform_bucket, - convert_parquet_files_to_dfs, - upload_dfs_to_database, -) +from src.load_lambda import * @pytest.fixture(scope="class") @@ -27,14 +21,43 @@ def aws_credentials(): def mock_s3_client(aws_credentials): with mock_aws(): yield boto3.client("s3") + +@pytest.fixture(scope="class") +def mock_sm_client(aws_credentials): + with mock_aws(): + yield boto3.client("secretsmanager") + + +@pytest.fixture +def mock_parquet_file(mocker): + return mocker.patch("src.load_lambda.convert_parquet_files_to_dfs") class TestLambdaHandler: pass class TestRetrieveSecrets: - pass + def test_retrieve_secrets_returns_dictionary(self, mock_sm_client): + secret = { + "cohort_id": "test_cohort_id", + "user": "test_user_id", + "password": "test_password", + "host": "test_host", + "database": "test_database", + "port": "test_port", + } + + secret_name = "test_secret" + + mock_sm_client.create_secret( + Name=secret_name, SecretString=json.dumps(secret) + ) + + result = retrieve_secrets(mock_sm_client, secret_name) + + assert isinstance(result, dict) + class TestConnectToDBAndReturnEngine: @@ -88,6 +111,5 @@ class TestConvertParquetToDfs: # assert "dim_staff" in result -@pytest.fixture -def mock_parquet_file(mocker): - return mocker.patch(src.load_lambda.convert_parquet_files_to_dfs()) +def mock_connect_db(mocker): + return mocker.patch("src.load_lambda.connect_to_db_and_return_engine") \ No newline at end of file -- cgit v1.2.3 From c7bc31ec5e3d838b3d48791ad13dd20600d7578f Mon Sep 17 00:00:00 2001 From: Ellie Date: Tue, 27 Aug 2024 14:14:43 +0100 Subject: add passing retrieve secrets tests --- tests/test_load_lambda.py | 23 ++++++++++++++++++----- 1 file changed, 18 insertions(+), 5 deletions(-) (limited to 'tests/test_load_lambda.py') diff --git a/tests/test_load_lambda.py b/tests/test_load_lambda.py index 3df94e4..9b0a271 100644 --- a/tests/test_load_lambda.py +++ b/tests/test_load_lambda.py @@ -3,6 +3,7 @@ import pyarrow.parquet as pq from io import BytesIO from moto import mock_aws import boto3 +import botocore.exceptions import os import pytest from src.load_lambda import * @@ -29,10 +30,6 @@ def mock_sm_client(aws_credentials): yield boto3.client("secretsmanager") -@pytest.fixture -def mock_parquet_file(mocker): - return mocker.patch("src.load_lambda.convert_parquet_files_to_dfs") - class TestLambdaHandler: pass @@ -58,6 +55,19 @@ class TestRetrieveSecrets: assert isinstance(result, dict) + def test_retrieve_secrets_returns_correct_keys_and_values(self, mock_sm_client): + secret_name = "test_secret" + + result = retrieve_secrets(mock_sm_client, secret_name) + + assert result["user"] == "test_user_id" + assert result["password"] == "test_password" + + def test_retrieve_secrets_returns_client_error_if_no_secret(self, mock_sm_client): + secret_name = "another_test_secret" + + with pytest.raises(botocore.exceptions.ClientError) as error: + retrieve_secrets(mock_sm_client, secret_name) class TestConnectToDBAndReturnEngine: @@ -112,4 +122,7 @@ class TestConvertParquetToDfs: def mock_connect_db(mocker): - return mocker.patch("src.load_lambda.connect_to_db_and_return_engine") \ No newline at end of file + return mocker.patch("src.load_lambda.connect_to_db_and_return_engine") + +class TestUploadDfsToDatabase: + pass \ No newline at end of file -- cgit v1.2.3 From f6584f5f52bc8731a2076e2d692faf28b107647d Mon Sep 17 00:00:00 2001 From: Alex Schofield Date: Tue, 27 Aug 2024 15:20:13 +0100 Subject: wip: add test for parquet file conversion --- tests/test_load_lambda.py | 59 ++++++++++++++++++++++++++++++++++++++++------- 1 file changed, 51 insertions(+), 8 deletions(-) (limited to 'tests/test_load_lambda.py') diff --git a/tests/test_load_lambda.py b/tests/test_load_lambda.py index 9b0a271..b5821a4 100644 --- a/tests/test_load_lambda.py +++ b/tests/test_load_lambda.py @@ -7,6 +7,7 @@ import botocore.exceptions import os import pytest from src.load_lambda import * +import tempfile @pytest.fixture(scope="class") @@ -22,7 +23,7 @@ def aws_credentials(): def mock_s3_client(aws_credentials): with mock_aws(): yield boto3.client("s3") - + @pytest.fixture(scope="class") def mock_sm_client(aws_credentials): @@ -30,6 +31,11 @@ def mock_sm_client(aws_credentials): yield boto3.client("secretsmanager") +@pytest.fixture(scope="class") +def mock_connect_db(mocker): + return mocker.patch("src.load_lambda.connect_to_db_and_return_engine") + + class TestLambdaHandler: pass @@ -47,9 +53,7 @@ class TestRetrieveSecrets: secret_name = "test_secret" - mock_sm_client.create_secret( - Name=secret_name, SecretString=json.dumps(secret) - ) + mock_sm_client.create_secret(Name=secret_name, SecretString=json.dumps(secret)) result = retrieve_secrets(mock_sm_client, secret_name) @@ -71,7 +75,17 @@ class TestRetrieveSecrets: class TestConnectToDBAndReturnEngine: - pass + def test_returns_unsuccessful_connection_when_wrong_credentials(self): + sm_secret = { + "host": "host", + "port": "port", + "user": "user", + "password": "password", + "database": "database", + } + + with pytest.raises(Exception): + connect_to_db_and_return_engine(json.dumps(sm_secret)) class TestGetTransformBucket: @@ -120,9 +134,38 @@ class TestConvertParquetToDfs: # result = convert_parquet_files_to_dfs(bucket_name="transform_bucket", client=mock_s3_client) # assert "dim_staff" in result + def test_function_returns_dictionary_with_file_key_and_dataframe( + self, mock_s3_client + ): + with tempfile.TemporaryDirectory() as tmp: + d = { + "test": ["Hello", "Bye"], + "design_id": ["Hello", "Bye"], + "design_name": ["Hello", "Bye"], + "file_name": ["Hello", "Bye"], + "file_location": ["Hello", "Bye"], + "Hello": ["Hello", "Bye"], + } + + test_df = pd.DataFrame(data=d) + + path = os.path.join(tmp, "test_parquet.parquet") + + test_df.to_parquet(path, engine="pyarrow") + + with open(path, "rb") as p: + mock_s3_client.put_object( + Bucket="transform_bucket", Key="test_parquet.parquet", Body=p.read() + ) + + result = convert_parquet_files_to_dfs( + bucket_name="transform_bucket", client=mock_s3_client + ) + + assert "test_parquet.parquet" in result + + pd.testing.assert_frame_equal(result["test_parquet.parquet"], test_df) -def mock_connect_db(mocker): - return mocker.patch("src.load_lambda.connect_to_db_and_return_engine") class TestUploadDfsToDatabase: - pass \ No newline at end of file + pass -- cgit v1.2.3 From f5bccf178ea1ebce213efd0518af63d74b00a11c Mon Sep 17 00:00:00 2001 From: Alex Schofield Date: Tue, 27 Aug 2024 15:34:35 +0100 Subject: test: add lambda_handler tests --- tests/test_load_lambda.py | 27 +++++++++++++++++++++------ 1 file changed, 21 insertions(+), 6 deletions(-) (limited to 'tests/test_load_lambda.py') diff --git a/tests/test_load_lambda.py b/tests/test_load_lambda.py index b5821a4..98ab36b 100644 --- a/tests/test_load_lambda.py +++ b/tests/test_load_lambda.py @@ -31,13 +31,28 @@ def mock_sm_client(aws_credentials): yield boto3.client("secretsmanager") -@pytest.fixture(scope="class") -def mock_connect_db(mocker): - return mocker.patch("src.load_lambda.connect_to_db_and_return_engine") - - class TestLambdaHandler: - pass + def test_lambda_handler_returns_success(self, mocker): + mocker.patch( + "src.load_lambda.upload_dfs_to_database", + return_value={"uploaded": ["table_one", "table_two"]}, + ) + result = lambda_handler(None, None) + assert result["statusCode"] == 200 + assert "table_one" in result["body"] + assert "table_two" in result["body"] + + def test_lambda_handler_does_not_upload_anything(self, mocker): + mocker.patch( + "src.load_lambda.upload_dfs_to_database", + return_value={"uploaded": []}, + ) + result = lambda_handler(None, None) + assert result["statusCode"] == 200 + assert "No dataframes were uploaded" in result["body"] + + def test_lambda_handler_returns_exception(self, mocker): + pass class TestRetrieveSecrets: -- cgit v1.2.3 From 843f11c302a2a9089c3726342cd1231015f074f7 Mon Sep 17 00:00:00 2001 From: Alex Schofield Date: Tue, 27 Aug 2024 15:36:12 +0100 Subject: docs: add comments for upload tests --- tests/test_load_lambda.py | 3 +++ 1 file changed, 3 insertions(+) (limited to 'tests/test_load_lambda.py') diff --git a/tests/test_load_lambda.py b/tests/test_load_lambda.py index 98ab36b..a29b75a 100644 --- a/tests/test_load_lambda.py +++ b/tests/test_load_lambda.py @@ -183,4 +183,7 @@ class TestConvertParquetToDfs: class TestUploadDfsToDatabase: + # Full success test + # Partial success test + # Failure test pass -- cgit v1.2.3 From cbfc98a9f43b5a0dae95337057c18c9dc2a298e3 Mon Sep 17 00:00:00 2001 From: Alex Schofield Date: Tue, 27 Aug 2024 16:00:29 +0100 Subject: wip: update TestLambdaHandler & lambda_handler function --- src/load_lambda.py | 19 +++++++++++-------- tests/test_load_lambda.py | 12 +++++++++--- 2 files changed, 20 insertions(+), 11 deletions(-) (limited to 'tests/test_load_lambda.py') diff --git a/src/load_lambda.py b/src/load_lambda.py index 11d1d70..39fa27d 100644 --- a/src/load_lambda.py +++ b/src/load_lambda.py @@ -23,18 +23,21 @@ logging.getLogger("botocore").setLevel(logging.INFO) def lambda_handler(event, context): try: uploaded_tables = upload_dfs_to_database() - if not uploaded_tables["uploaded"]: + if uploaded_tables["not_uploaded"]: return { "statusCode": 200, "body": json.dumps("No dataframes were uploaded."), } - return { - "statusCode": 200, - "body": json.dumps( - f"""The following dataframes were uploaded successfully: - {uploaded_tables["uploaded"]} .""" - ), - } + + if uploaded_tables["uploaded"]: + return { + "statusCode": 200, + "body": json.dumps( + f"""The following dataframes were uploaded successfully: + {uploaded_tables["uploaded"]} .""" + ), + } + except Exception as e: logger.error(f"Error: {e}", exc_info=True) return {"statusCode": 500, "body": json.dumps("Internal server error.")} diff --git a/tests/test_load_lambda.py b/tests/test_load_lambda.py index a29b75a..9286e48 100644 --- a/tests/test_load_lambda.py +++ b/tests/test_load_lambda.py @@ -35,7 +35,7 @@ class TestLambdaHandler: def test_lambda_handler_returns_success(self, mocker): mocker.patch( "src.load_lambda.upload_dfs_to_database", - return_value={"uploaded": ["table_one", "table_two"]}, + return_value={"uploaded": ["table_one", "table_two"], "not_uploaded": []}, ) result = lambda_handler(None, None) assert result["statusCode"] == 200 @@ -45,14 +45,20 @@ class TestLambdaHandler: def test_lambda_handler_does_not_upload_anything(self, mocker): mocker.patch( "src.load_lambda.upload_dfs_to_database", - return_value={"uploaded": []}, + return_value={"uploaded": [], "not_uploaded": []}, ) result = lambda_handler(None, None) assert result["statusCode"] == 200 assert "No dataframes were uploaded" in result["body"] def test_lambda_handler_returns_exception(self, mocker): - pass + mocker.patch( + "src.load_lambda.upload_dfs_to_database", + return_value={"test": []}, + ) + + with pytest.raises(Exception): + lambda_handler(None, None) class TestRetrieveSecrets: -- cgit v1.2.3 From 0ea88c0216d9e5eca9e4aca4f2fa427d38184648 Mon Sep 17 00:00:00 2001 From: Ellie Date: Tue, 27 Aug 2024 16:40:21 +0100 Subject: add passing tests for lambda handler --- tests/test_load_lambda.py | 16 +++++++++------- 1 file changed, 9 insertions(+), 7 deletions(-) (limited to 'tests/test_load_lambda.py') diff --git a/tests/test_load_lambda.py b/tests/test_load_lambda.py index 9286e48..0b13b54 100644 --- a/tests/test_load_lambda.py +++ b/tests/test_load_lambda.py @@ -32,7 +32,7 @@ def mock_sm_client(aws_credentials): class TestLambdaHandler: - def test_lambda_handler_returns_success(self, mocker): + def test_lambda_handler_returns_200_and_table_name_if_uploaded(self, mocker): mocker.patch( "src.load_lambda.upload_dfs_to_database", return_value={"uploaded": ["table_one", "table_two"], "not_uploaded": []}, @@ -42,23 +42,25 @@ class TestLambdaHandler: assert "table_one" in result["body"] assert "table_two" in result["body"] - def test_lambda_handler_does_not_upload_anything(self, mocker): + def test_lambda_handler_returns_200_and_table_name_if_not_uploaded(self, mocker): mocker.patch( "src.load_lambda.upload_dfs_to_database", - return_value={"uploaded": [], "not_uploaded": []}, + return_value={"uploaded": [], "not_uploaded": ["table_one"]}, ) result = lambda_handler(None, None) assert result["statusCode"] == 200 assert "No dataframes were uploaded" in result["body"] - def test_lambda_handler_returns_exception(self, mocker): + def test_lambda_handler_returns_error_if_both_lists_empty(self, mocker): mocker.patch( "src.load_lambda.upload_dfs_to_database", - return_value={"test": []}, + return_value={"uploaded": [], "not_uploaded": []}, ) - with pytest.raises(Exception): - lambda_handler(None, None) + result = lambda_handler(None, None) + + assert result == {"error"} + class TestRetrieveSecrets: -- cgit v1.2.3 From 57617571df0a667aca55fc54184696a19c689524 Mon Sep 17 00:00:00 2001 From: Ellie Date: Tue, 27 Aug 2024 17:00:08 +0100 Subject: add lambda handler updated tests --- tests/test_load_lambda.py | 1 + 1 file changed, 1 insertion(+) (limited to 'tests/test_load_lambda.py') diff --git a/tests/test_load_lambda.py b/tests/test_load_lambda.py index 0b13b54..829b908 100644 --- a/tests/test_load_lambda.py +++ b/tests/test_load_lambda.py @@ -63,6 +63,7 @@ class TestLambdaHandler: + class TestRetrieveSecrets: def test_retrieve_secrets_returns_dictionary(self, mock_sm_client): secret = { -- cgit v1.2.3 From 08c971f0e56d0896aa09200c26b5cfa53ff29ca1 Mon Sep 17 00:00:00 2001 From: Ellie Date: Tue, 27 Aug 2024 17:27:40 +0100 Subject: add json.loads to retrieve secrets --- tests/test_load_lambda.py | 4 ++-- 1 file changed, 2 insertions(+), 2 deletions(-) (limited to 'tests/test_load_lambda.py') diff --git a/tests/test_load_lambda.py b/tests/test_load_lambda.py index 829b908..02cf2c0 100644 --- a/tests/test_load_lambda.py +++ b/tests/test_load_lambda.py @@ -79,14 +79,14 @@ class TestRetrieveSecrets: mock_sm_client.create_secret(Name=secret_name, SecretString=json.dumps(secret)) - result = retrieve_secrets(mock_sm_client, secret_name) + result = json.loads(retrieve_secrets(mock_sm_client, secret_name)) assert isinstance(result, dict) def test_retrieve_secrets_returns_correct_keys_and_values(self, mock_sm_client): secret_name = "test_secret" - result = retrieve_secrets(mock_sm_client, secret_name) + result = json.loads(retrieve_secrets(mock_sm_client, secret_name)) assert result["user"] == "test_user_id" assert result["password"] == "test_password" -- cgit v1.2.3 From 95935534931b5ff6e617ba74c86cb7a6718128e4 Mon Sep 17 00:00:00 2001 From: "deepsource-autofix[bot]" <62050782+deepsource-autofix[bot]@users.noreply.github.com> Date: Wed, 28 Aug 2024 08:24:21 +0000 Subject: style: format code with Autopep8, Black and Ruff Formatter This commit fixes the style issues introduced in 08c971f according to the output from Autopep8, Black and Ruff Formatter. Details: https://github.com/ajschofield/de-project-bentley/pull/102 --- src/dataframes.py | 182 ++++++++++++++++++++++++---------------------- tests/test_dataframes.py | 43 ++++++----- tests/test_load_lambda.py | 2 - 3 files changed, 123 insertions(+), 104 deletions(-) (limited to 'tests/test_load_lambda.py') diff --git a/src/dataframes.py b/src/dataframes.py index 4b32b36..43facd6 100644 --- a/src/dataframes.py +++ b/src/dataframes.py @@ -20,8 +20,11 @@ import requests def create_fact_sales_order(dict_of_df): df_sales = dict_of_df["sales_order"] df_sales.index.name = "sales_record_id" -<<<<<<< HEAD - df_sales["created_date"] = df_sales["created_at"].astype("datetime64[ns]").dt.date + + +<< << << < HEAD + df_sales["created_date"] = df_sales["created_at"].astype( + "datetime64[ns]").dt.date df_sales["created_time"] = ( df_sales["created_at"].astype("datetime64[ns]").dt.floor("s").dt.time ) @@ -30,27 +33,29 @@ def create_fact_sales_order(dict_of_df): ) df_sales["last_updated_time"] = ( df_sales["last_updated"].astype("datetime64[ns]").dt.floor("s").dt.time -======= - df_sales["created_date"] = pd.to_datetime(df_sales["created_at"], format="%Y-%m-%d") - df_sales["created_time"] = pd.to_datetime(df_sales["created_at"], format="%H-%M-%S") - df_sales["last_updated_date"] = pd.to_datetime( +== == == = + df_sales["created_date"]=pd.to_datetime( + df_sales["created_at"], format="%Y-%m-%d") + df_sales["created_time"]=pd.to_datetime( + df_sales["created_at"], format="%H-%M-%S") + df_sales["last_updated_date"]=pd.to_datetime( df_sales["last_updated"], format="%Y-%m-%d" ) - df_sales["last_updated_time"] = pd.to_datetime( + df_sales["last_updated_time"]=pd.to_datetime( df_sales["last_updated"], format="%H-%M-%S" ->>>>>>> 5db3f61 (style: format code with Autopep8, Black and Ruff Formatter) +>> >>>> > 5db3f61(style: format code with Autopep8, Black and Ruff Formatter) ) - df_sales["agreed_delivery_date"] = pd.to_datetime( + df_sales["agreed_delivery_date"]=pd.to_datetime( df_sales["agreed_delivery_date"], format="%Y-%m-%d" ) - df_sales["agreed_payment_date"] = pd.to_datetime( + df_sales["agreed_payment_date"]=pd.to_datetime( df_sales["agreed_payment_date"], format="%Y-%m-%d" ) -<<<<<<< HEAD - df_sales = df_sales.drop(labels=["created_at", "last_updated"], axis=1) -======= +<< << << < HEAD + df_sales=df_sales.drop(labels=["created_at", "last_updated"], axis=1) +== == == = df_sales.drop(labels=["created_at", "last_updated"], axis=1, inplace=True) ->>>>>>> 5db3f61 (style: format code with Autopep8, Black and Ruff Formatter) +>> >>>> > 5db3f61(style: format code with Autopep8, Black and Ruff Formatter) df_sales.reset_index(inplace=True) return df_sales @@ -59,37 +64,40 @@ def create_fact_sales_order(dict_of_df): def create_fact_purchase_orders(dict_of_df): - df_po = dict_of_df["purchase_order"] - df_po.index.name = "purchase_record_id" -<<<<<<< HEAD - df_po["created_date"] = df_po["created_at"].astype("datetime64[ns]").dt.date - df_po["created_time"] = ( + df_po=dict_of_df["purchase_order"] + df_po.index.name="purchase_record_id" +<< << << < HEAD + df_po["created_date"]=df_po["created_at"].astype("datetime64[ns]").dt.date + df_po["created_time"]=( df_po["created_at"].astype("datetime64[ns]").dt.floor("s").dt.time ) - df_po["last_updated_date"] = df_po["last_updated"].astype("datetime64[ns]").dt.date - df_po["last_updated_time"] = ( + df_po["last_updated_date"]=df_po["last_updated"].astype( + "datetime64[ns]").dt.date + df_po["last_updated_time"]=( df_po["last_updated"].astype("datetime64[ns]").dt.floor("s").dt.time -======= - df_po["created_date"] = pd.to_datetime(df_po["created_at"], format="%Y-%m-%d") - df_po["created_time"] = pd.to_datetime(df_po["created_at"], format="%H-%M-%S") - df_po["last_updated_date"] = pd.to_datetime( +== == == = + df_po["created_date"]=pd.to_datetime( + df_po["created_at"], format="%Y-%m-%d") + df_po["created_time"]=pd.to_datetime( + df_po["created_at"], format="%H-%M-%S") + df_po["last_updated_date"]=pd.to_datetime( df_po["last_updated"], format="%Y-%m-%d" ) - df_po["last_updated_time"] = pd.to_datetime( + df_po["last_updated_time"]=pd.to_datetime( df_po["last_updated"], format="%H-%M-%S" ->>>>>>> 5db3f61 (style: format code with Autopep8, Black and Ruff Formatter) +>> >>>> > 5db3f61(style: format code with Autopep8, Black and Ruff Formatter) ) - df_po["agreed_delivery_date"] = pd.to_datetime( + df_po["agreed_delivery_date"]=pd.to_datetime( df_po["agreed_delivery_date"], format="%Y-%m-%d" ) - df_po["agreed_payment_date"] = pd.to_datetime( + df_po["agreed_payment_date"]=pd.to_datetime( df_po["agreed_payment_date"], format="%Y-%m-%d" ) -<<<<<<< HEAD - df_po = df_po.drop(labels=["created_at", "last_updated"], axis=1) -======= +<< << << < HEAD + df_po=df_po.drop(labels=["created_at", "last_updated"], axis=1) +== == == = df_po.drop(labels=["created_at", "last_updated"], axis=1, inplace=True) ->>>>>>> 5db3f61 (style: format code with Autopep8, Black and Ruff Formatter) +>> >>>> > 5db3f61(style: format code with Autopep8, Black and Ruff Formatter) df_po.reset_index(inplace=True) return df_po @@ -98,42 +106,44 @@ def create_fact_purchase_orders(dict_of_df): def create_fact_payment(dict_of_df): - df_payment = dict_of_df["payment"] - df_payment.index.name = "payment_record_id" -<<<<<<< HEAD - df_payment["created_date"] = ( + df_payment=dict_of_df["payment"] + df_payment.index.name="payment_record_id" +<< << << < HEAD + df_payment["created_date"]=( df_payment["created_at"].astype("datetime64[ns]").dt.date ) - df_payment["created_time"] = ( + df_payment["created_time"]=( df_payment["created_at"].astype("datetime64[ns]").dt.floor("s").dt.time ) - df_payment["last_updated_date"] = ( + df_payment["last_updated_date"]=( df_payment["last_updated"].astype("datetime64[ns]").dt.date ) - df_payment["last_updated_time"] = ( - df_payment["last_updated"].astype("datetime64[ns]").dt.floor("s").dt.time -======= - df_payment["created_date"] = pd.to_datetime( + df_payment["last_updated_time"]=( + df_payment["last_updated"].astype( + "datetime64[ns]").dt.floor("s").dt.time +== == == = + df_payment["created_date"]=pd.to_datetime( df_payment["created_at"], format="%Y-%m-%d" ) - df_payment["created_time"] = pd.to_datetime( + df_payment["created_time"]=pd.to_datetime( df_payment["created_at"], format="%H-%M-%S" ) - df_payment["last_updated_date"] = pd.to_datetime( + df_payment["last_updated_date"]=pd.to_datetime( df_payment["last_updated"], format="%Y-%m-%d" ) - df_payment["last_updated_time"] = pd.to_datetime( + df_payment["last_updated_time"]=pd.to_datetime( df_payment["last_updated"], format="%H-%M-%S" ->>>>>>> 5db3f61 (style: format code with Autopep8, Black and Ruff Formatter) +>> >>>> > 5db3f61(style: format code with Autopep8, Black and Ruff Formatter) ) - df_payment["payment_date"] = pd.to_datetime( + df_payment["payment_date"]=pd.to_datetime( df_payment["payment_date"], format="%Y-%m-%d" ) -<<<<<<< HEAD - df_payment = df_payment.drop(labels=["created_at", "last_updated"], axis=1) -======= - df_payment.drop(labels=["created_at", "last_updated"], axis=1, inplace=True) ->>>>>>> 5db3f61 (style: format code with Autopep8, Black and Ruff Formatter) +<< << << < HEAD + df_payment=df_payment.drop(labels=["created_at", "last_updated"], axis=1) +== == == = + df_payment.drop( + labels=["created_at", "last_updated"], axis=1, inplace=True) +>> >>>> > 5db3f61(style: format code with Autopep8, Black and Ruff Formatter) df_payment.reset_index(inplace=True) return df_payment @@ -142,7 +152,7 @@ def create_fact_payment(dict_of_df): def create_dim_transaction(dict_of_df): - df_transaction = dict_of_df["transaction"].drop( + df_transaction=dict_of_df["transaction"].drop( labels=["created_at", "last_updated"], axis=1 ) return df_transaction @@ -152,7 +162,7 @@ def create_dim_transaction(dict_of_df): def create_dim_location(dict_of_df): - df_loc = ( + df_loc=( dict_of_df["address"] .drop(labels=["created_at", "last_updated"], axis=1) .rename(columns={"address_id": "location_id"}) @@ -161,10 +171,10 @@ def create_dim_location(dict_of_df): def create_dim_counterparty(dict_of_df): - df_prefixed_address = dict_of_df["address"].drop(labels=["created_at", "last_updated"], axis=1).add_prefix( + df_prefixed_address=dict_of_df["address"].drop(labels=["created_at", "last_updated"], axis=1).add_prefix( "counterparty_legal_", axis=1 ) - df_cp = pd.merge( + df_cp=pd.merge( dict_of_df["counterparty"], df_prefixed_address, left_on="legal_address_id", @@ -181,32 +191,32 @@ def create_dim_counterparty(dict_of_df): def create_dim_date(dict_of_df): - fact_dfs = [ + fact_dfs=[ create_fact_payment(dict_of_df), create_fact_purchase_orders(dict_of_df), create_fact_sales_order(dict_of_df), ] - list_of_date_columns = [] + list_of_date_columns=[] for df in fact_dfs: - date_col_names = [ -<<<<<<< HEAD + date_col_names=[ +<< << << < HEAD col_name for col_name in list(df.columns) if "_date" in col_name -======= +== == == = col_name for col_name in list(df.columns) if "date" in col_name ->>>>>>> 5db3f61 (style: format code with Autopep8, Black and Ruff Formatter) +>> >>>> > 5db3f61(style: format code with Autopep8, Black and Ruff Formatter) ] for col in date_col_names: list_of_date_columns.append(df[col]) - sr_date = pd.array(pd.concat(list_of_date_columns), dtype="datetime64[ns]") - df_date = pd.DataFrame(data=sr_date, columns=["date_id"]) + sr_date=pd.array(pd.concat(list_of_date_columns), dtype="datetime64[ns]") + df_date=pd.DataFrame(data=sr_date, columns=["date_id"]) df_date.drop_duplicates(inplace=True) - df_date["year"] = df_date["date_id"].dt.year - df_date["month"] = df_date["date_id"].dt.month - df_date["day"] = df_date["date_id"].dt.day - df_date["day_of_week"] = df_date["date_id"].dt.dayofweek - df_date["day_name"] = df_date["date_id"].dt.day_name() - df_date["month_name"] = df_date["date_id"].dt.month_name() - df_date["quarter"] = df_date["date_id"].dt.quarter + df_date["year"]=df_date["date_id"].dt.year + df_date["month"]=df_date["date_id"].dt.month + df_date["day"]=df_date["date_id"].dt.day + df_date["day_of_week"]=df_date["date_id"].dt.dayofweek + df_date["day_name"]=df_date["date_id"].dt.day_name() + df_date["month_name"]=df_date["date_id"].dt.month_name() + df_date["quarter"]=df_date["date_id"].dt.quarter return df_date @@ -214,13 +224,13 @@ def create_dim_date(dict_of_df): def scrape_currency_names(): - response = requests.get("https://www.xe.com/currency/").content - soup = BeautifulSoup(response, "html.parser") - currency = [ + response=requests.get("https://www.xe.com/currency/").content + soup=BeautifulSoup(response, "html.parser") + currency=[ item.text for item in soup.findAll("a", attrs={"class": "sc-299dec64-6 fZPTSw"}) ] - sr = pd.Series(currency) - df_cur = sr.str.split(pat=" - ", expand=True).rename( + sr=pd.Series(currency) + df_cur=sr.str.split(pat=" - ", expand=True).rename( {0: "currency_code", 1: "currency_name"}, axis=1 ) return df_cur @@ -230,8 +240,9 @@ def scrape_currency_names(): def create_dim_currency(dict_of_df, names=scrape_currency_names()): - df_cur = dict_of_df["currency"].drop(labels=["created_at", "last_updated"], axis=1) - dim_cur = pd.merge( + df_cur=dict_of_df["currency"].drop( + labels=["created_at", "last_updated"], axis=1) + dim_cur=pd.merge( df_cur, names, left_on="currency_code", right_on="currency_code", how="inner" ) return dim_cur @@ -241,8 +252,9 @@ def create_dim_currency(dict_of_df, names=scrape_currency_names()): def create_dim_payment_type(dict_of_df): - df_payment_type = dict_of_df["payment_type"] - dim_payment_type = df_payment_type.loc[:, ["payment_type_id", "payment_type_name"]] + df_payment_type=dict_of_df["payment_type"] + dim_payment_type=df_payment_type.loc[:, [ + "payment_type_id", "payment_type_name"]] return dim_payment_type @@ -250,8 +262,8 @@ def create_dim_payment_type(dict_of_df): def create_dim_design(dict_of_df): - df_design = dict_of_df["design"] - dim_design = df_design.loc[ + df_design=dict_of_df["design"] + dim_design=df_design.loc[ :, ["design_id", "design_name", "file_name", "file_location"] ] return dim_design @@ -261,10 +273,10 @@ def create_dim_design(dict_of_df): def create_dim_staff(dict_of_df): - staff_department = pd.merge( + staff_department=pd.merge( dict_of_df["staff"], dict_of_df["department"], on="department_id", how="left" ) - dim_staff = staff_department.loc[ + dim_staff=staff_department.loc[ :, [ "staff_id", diff --git a/tests/test_dataframes.py b/tests/test_dataframes.py index cc133fe..785a3fd 100644 --- a/tests/test_dataframes.py +++ b/tests/test_dataframes.py @@ -54,7 +54,8 @@ class TestCreateDimStaff: "email_address": ["Hello", "Bye"], "department_id": ["Hello", "Bye"], } - test_df = {"staff": pd.DataFrame(data=d), "department": pd.DataFrame(data=d2)} + test_df = {"staff": pd.DataFrame( + data=d), "department": pd.DataFrame(data=d2)} result = create_dim_staff(test_df) assert isinstance(result, pd.DataFrame) @@ -71,7 +72,8 @@ class TestCreateDimStaff: "email_address": ["Hello", "Bye"], "department_id": ["Hello", "Bye"], } - test_df = {"staff": pd.DataFrame(data=d), "department": pd.DataFrame(data=d2)} + test_df = {"staff": pd.DataFrame( + data=d), "department": pd.DataFrame(data=d2)} result = create_dim_staff(test_df) expected_d = { "staff_id": ["Hello", "Bye"], @@ -88,7 +90,8 @@ class TestCreateDimStaff: class TestCreatePaymentType: def test_create_dim_payment_type_returns_correct_columns_and_values(self): - d = {"payment_type_id": ["Hello", "Bye"], "payment_type_name": ["Hello", "Bye"]} + d = {"payment_type_id": ["Hello", "Bye"], + "payment_type_name": ["Hello", "Bye"]} test_df = {"payment_type": pd.DataFrame(data=d)} result = create_dim_payment_type(test_df) expected_columns = ["payment_type_id", "payment_type_name"] @@ -180,11 +183,13 @@ class TestCreateDimDate: index=[0], ) df_two = pd.DataFrame( - data={"updated_date": dt(2020, 5, 17), "created_date": dt(2021, 9, 13)}, + data={"updated_date": dt(2020, 5, 17), + "created_date": dt(2021, 9, 13)}, index=[0], ) df_three = pd.DataFrame( - data={"updated_date": dt(2022, 5, 17), "created_date": dt(2023, 5, 13)}, + data={"updated_date": dt(2022, 5, 17), + "created_date": dt(2023, 5, 13)}, index=[0], ) expected_df = pd.DataFrame( @@ -214,7 +219,8 @@ class TestCreateDimDate: mock_fso.return_value = df_three result = create_dim_date({"dum": 0}) result.reset_index(inplace=True, drop=True) - assert result.eq(expected_df, axis="columns").all(axis=None) + assert result.eq( + expected_df, axis="columns").all(axis=None) class TestCreateDimLocation: @@ -222,7 +228,8 @@ class TestCreateDimLocation: dict_df = { "address": pd.DataFrame( data=[["some_time", "some_other_time", 1, "SE18 9QO"]], - columns=["created_at", "last_updated", "address_id", "postal_code"], + columns=["created_at", "last_updated", + "address_id", "postal_code"], ) } result = create_dim_location(dict_df) @@ -252,7 +259,7 @@ class TestCreateFactPayment: "payment": pd.DataFrame( data=[ [ -<<<<<<< HEAD + << << << < HEAD dt.strptime( "2022-11-03 14:20:49.962846", "%Y-%m-%d %H:%M:%S.%f" ), @@ -262,13 +269,13 @@ class TestCreateFactPayment: 1, "SE18 9QO", "2020-07-16", -======= + == == === dt(2020, 5, 17, 6, 15, 20), dt(2020, 5, 20, 8, 19, 30), 1, "SE18 9QO", "2020-7-16", ->>>>>>> 5db3f61 (style: format code with Autopep8, Black and Ruff Formatter) + >>>>>> > 5db3f61(style: format code with Autopep8, Black and Ruff Formatter) ] ], columns=[ @@ -295,10 +302,12 @@ class TestCreateFactPayment: for col in list(result.columns): assert col in expected_cols for col in expected_cols: -<<<<<<< HEAD - if "_date" or "_time" in col: - assert result[col].dtype == "O" -======= - if "date" in col: - assert result[col].dtype == "datetime64[ns]" ->>>>>>> 5db3f61 (style: format code with Autopep8, Black and Ruff Formatter) + + +<< << << < HEAD +if "_date" or "_time" in col: + assert result[col].dtype == "O" +== == == = +if "date" in col: + assert result[col].dtype == "datetime64[ns]" +>>>>>> > 5db3f61(style: format code with Autopep8, Black and Ruff Formatter) diff --git a/tests/test_load_lambda.py b/tests/test_load_lambda.py index 02cf2c0..65106f7 100644 --- a/tests/test_load_lambda.py +++ b/tests/test_load_lambda.py @@ -62,8 +62,6 @@ class TestLambdaHandler: assert result == {"error"} - - class TestRetrieveSecrets: def test_retrieve_secrets_returns_dictionary(self, mock_sm_client): secret = { -- cgit v1.2.3