# Remove duplicates, keeping the first instance df.drop_duplicates(inplace=True) print(f”Dataframe size after removing duplicates: {len(df)} rows.”)